Chris@23: /* -*- c-basic-offset: 4 indent-tabs-mode: nil -*- vi:set ts=8 sts=4 sw=4: */ Chris@23: Chris@23: /* Chris@23: Vamp Chris@23: Chris@23: An API for audio analysis and feature extraction plugins. Chris@23: Chris@23: Centre for Digital Music, Queen Mary, University of London. Chris@23: Copyright 2006 Chris Cannam. Chris@23: Chris@23: Permission is hereby granted, free of charge, to any person Chris@23: obtaining a copy of this software and associated documentation Chris@23: files (the "Software"), to deal in the Software without Chris@23: restriction, including without limitation the rights to use, copy, Chris@23: modify, merge, publish, distribute, sublicense, and/or sell copies Chris@23: of the Software, and to permit persons to whom the Software is Chris@23: furnished to do so, subject to the following conditions: Chris@23: Chris@23: The above copyright notice and this permission notice shall be Chris@23: included in all copies or substantial portions of the Software. Chris@23: Chris@23: THE SOFTWARE IS PROVIDED "AS IS", WITHOUT WARRANTY OF ANY KIND, Chris@23: EXPRESS OR IMPLIED, INCLUDING BUT NOT LIMITED TO THE WARRANTIES OF Chris@23: MERCHANTABILITY, FITNESS FOR A PARTICULAR PURPOSE AND Chris@23: NONINFRINGEMENT. IN NO EVENT SHALL THE AUTHORS BE LIABLE FOR Chris@23: ANY CLAIM, DAMAGES OR OTHER LIABILITY, WHETHER IN AN ACTION OF Chris@23: CONTRACT, TORT OR OTHERWISE, ARISING FROM, OUT OF OR IN CONNECTION Chris@23: WITH THE SOFTWARE OR THE USE OR OTHER DEALINGS IN THE SOFTWARE. Chris@23: Chris@23: Except as contained in this notice, the names of the Centre for Chris@23: Digital Music; Queen Mary, University of London; and Chris Cannam Chris@23: shall not be used in advertising or otherwise to promote the sale, Chris@23: use or other dealings in this Software without prior written Chris@23: authorization. Chris@23: */ Chris@23: Chris@23: #include "PercussionOnsetDetector.h" Chris@23: Chris@23: using std::string; Chris@23: using std::vector; Chris@23: using std::cerr; Chris@23: using std::endl; Chris@23: Chris@23: #include Chris@23: Chris@23: Chris@23: PercussionOnsetDetector::PercussionOnsetDetector(float inputSampleRate) : Chris@23: Plugin(inputSampleRate), Chris@23: m_stepSize(0), Chris@23: m_blockSize(0), Chris@23: m_threshold(3), Chris@23: m_sensitivity(40), Chris@23: m_priorMagnitudes(0), Chris@23: m_dfMinus1(0), Chris@23: m_dfMinus2(0) Chris@23: { Chris@23: } Chris@23: Chris@23: PercussionOnsetDetector::~PercussionOnsetDetector() Chris@23: { Chris@23: delete[] m_priorMagnitudes; Chris@23: } Chris@23: Chris@23: string Chris@23: PercussionOnsetDetector::getIdentifier() const Chris@23: { Chris@23: return "percussiononsets"; Chris@23: } Chris@23: Chris@23: string Chris@23: PercussionOnsetDetector::getName() const Chris@23: { Chris@23: return "Simple Percussion Onset Detector"; Chris@23: } Chris@23: Chris@23: string Chris@23: PercussionOnsetDetector::getDescription() const Chris@23: { Chris@23: return "Detect percussive note onsets by identifying broadband energy rises"; Chris@23: } Chris@23: Chris@23: string Chris@23: PercussionOnsetDetector::getMaker() const Chris@23: { Chris@23: return "Vamp SDK Example Plugins"; Chris@23: } Chris@23: Chris@23: int Chris@23: PercussionOnsetDetector::getPluginVersion() const Chris@23: { Chris@23: return 2; Chris@23: } Chris@23: Chris@23: string Chris@23: PercussionOnsetDetector::getCopyright() const Chris@23: { Chris@23: return "Code copyright 2006 Queen Mary, University of London, after Dan Barry et al 2005. Freely redistributable (BSD license)"; Chris@23: } Chris@23: Chris@23: size_t Chris@23: PercussionOnsetDetector::getPreferredStepSize() const Chris@23: { Chris@23: return 0; Chris@23: } Chris@23: Chris@23: size_t Chris@23: PercussionOnsetDetector::getPreferredBlockSize() const Chris@23: { Chris@23: return 1024; Chris@23: } Chris@23: Chris@23: bool Chris@23: PercussionOnsetDetector::initialise(size_t channels, size_t stepSize, size_t blockSize) Chris@23: { Chris@23: if (channels < getMinChannelCount() || Chris@23: channels > getMaxChannelCount()) return false; Chris@23: Chris@23: m_stepSize = stepSize; Chris@23: m_blockSize = blockSize; Chris@23: Chris@23: m_priorMagnitudes = new float[m_blockSize/2]; Chris@23: Chris@23: for (size_t i = 0; i < m_blockSize/2; ++i) { Chris@23: m_priorMagnitudes[i] = 0.f; Chris@23: } Chris@23: Chris@23: m_dfMinus1 = 0.f; Chris@23: m_dfMinus2 = 0.f; Chris@23: Chris@23: return true; Chris@23: } Chris@23: Chris@23: void Chris@23: PercussionOnsetDetector::reset() Chris@23: { Chris@23: for (size_t i = 0; i < m_blockSize/2; ++i) { Chris@23: m_priorMagnitudes[i] = 0.f; Chris@23: } Chris@23: Chris@23: m_dfMinus1 = 0.f; Chris@23: m_dfMinus2 = 0.f; Chris@23: } Chris@23: Chris@23: PercussionOnsetDetector::ParameterList Chris@23: PercussionOnsetDetector::getParameterDescriptors() const Chris@23: { Chris@23: ParameterList list; Chris@23: Chris@23: ParameterDescriptor d; Chris@23: d.identifier = "threshold"; Chris@23: d.name = "Energy rise threshold"; Chris@23: d.description = "Energy rise within a frequency bin necessary to count toward broadband total"; Chris@23: d.unit = "dB"; Chris@23: d.minValue = 0; Chris@23: d.maxValue = 20; Chris@23: d.defaultValue = 3; Chris@23: d.isQuantized = false; Chris@23: list.push_back(d); Chris@23: Chris@23: d.identifier = "sensitivity"; Chris@23: d.name = "Sensitivity"; Chris@23: d.description = "Sensitivity of peak detector applied to broadband detection function"; Chris@23: d.unit = "%"; Chris@23: d.minValue = 0; Chris@23: d.maxValue = 100; Chris@23: d.defaultValue = 40; Chris@23: d.isQuantized = false; Chris@23: list.push_back(d); Chris@23: Chris@23: return list; Chris@23: } Chris@23: Chris@23: float Chris@23: PercussionOnsetDetector::getParameter(std::string id) const Chris@23: { Chris@23: if (id == "threshold") return m_threshold; Chris@23: if (id == "sensitivity") return m_sensitivity; Chris@23: return 0.f; Chris@23: } Chris@23: Chris@23: void Chris@23: PercussionOnsetDetector::setParameter(std::string id, float value) Chris@23: { Chris@23: if (id == "threshold") { Chris@23: if (value < 0) value = 0; Chris@23: if (value > 20) value = 20; Chris@23: m_threshold = value; Chris@23: } else if (id == "sensitivity") { Chris@23: if (value < 0) value = 0; Chris@23: if (value > 100) value = 100; Chris@23: m_sensitivity = value; Chris@23: } Chris@23: } Chris@23: Chris@23: PercussionOnsetDetector::OutputList Chris@23: PercussionOnsetDetector::getOutputDescriptors() const Chris@23: { Chris@23: OutputList list; Chris@23: Chris@23: OutputDescriptor d; Chris@23: d.identifier = "onsets"; Chris@23: d.name = "Onsets"; Chris@23: d.description = "Percussive note onset locations"; Chris@23: d.unit = ""; Chris@23: d.hasFixedBinCount = true; Chris@23: d.binCount = 0; Chris@23: d.hasKnownExtents = false; Chris@23: d.isQuantized = false; Chris@23: d.sampleType = OutputDescriptor::VariableSampleRate; Chris@23: d.sampleRate = m_inputSampleRate; Chris@23: list.push_back(d); Chris@23: Chris@23: d.identifier = "detectionfunction"; Chris@23: d.name = "Detection Function"; Chris@23: d.description = "Broadband energy rise detection function"; Chris@23: d.binCount = 1; Chris@23: d.isQuantized = true; Chris@23: d.quantizeStep = 1.0; Chris@23: d.sampleType = OutputDescriptor::OneSamplePerStep; Chris@23: list.push_back(d); Chris@23: Chris@23: return list; Chris@23: } Chris@23: Chris@23: PercussionOnsetDetector::FeatureSet Chris@23: PercussionOnsetDetector::process(const float *const *inputBuffers, Chris@23: Vamp::RealTime ts) Chris@23: { Chris@23: if (m_stepSize == 0) { Chris@23: cerr << "ERROR: PercussionOnsetDetector::process: " Chris@23: << "PercussionOnsetDetector has not been initialised" Chris@23: << endl; Chris@23: return FeatureSet(); Chris@23: } Chris@23: Chris@23: int count = 0; Chris@23: Chris@23: for (size_t i = 1; i < m_blockSize/2; ++i) { Chris@23: Chris@23: float real = inputBuffers[0][i*2]; Chris@23: float imag = inputBuffers[0][i*2 + 1]; Chris@23: Chris@23: float sqrmag = real * real + imag * imag; Chris@23: Chris@23: if (m_priorMagnitudes[i] > 0.f) { Chris@23: float diff = 10.f * log10f(sqrmag / m_priorMagnitudes[i]); Chris@23: Chris@23: // std::cout << "i=" << i << ", sqrmag=" << sqrmag << ", prior=" << m_priorMagnitudes[i] << ", diff=" << diff << ", threshold=" << m_threshold << " " << (diff >= m_threshold ? "[*]" : "") << std::endl; Chris@23: Chris@23: if (diff >= m_threshold) ++count; Chris@23: } Chris@23: Chris@23: m_priorMagnitudes[i] = sqrmag; Chris@23: } Chris@23: Chris@23: FeatureSet returnFeatures; Chris@23: Chris@23: Feature detectionFunction; Chris@23: detectionFunction.hasTimestamp = false; Chris@23: detectionFunction.values.push_back(count); Chris@23: returnFeatures[1].push_back(detectionFunction); Chris@23: Chris@23: if (m_dfMinus2 < m_dfMinus1 && Chris@23: m_dfMinus1 >= count && Chris@23: m_dfMinus1 > ((100 - m_sensitivity) * m_blockSize) / 200) { Chris@23: Chris@23: //std::cout << "result at " << ts << "! (count == " << count << ", prev == " << m_dfMinus1 << ")" << std::endl; Chris@23: Chris@23: Feature onset; Chris@23: onset.hasTimestamp = true; Chris@23: onset.timestamp = ts - Vamp::RealTime::frame2RealTime Chris@23: (m_stepSize, int(m_inputSampleRate + 0.5)); Chris@23: returnFeatures[0].push_back(onset); Chris@23: } Chris@23: Chris@23: m_dfMinus2 = m_dfMinus1; Chris@23: m_dfMinus1 = count; Chris@23: Chris@23: return returnFeatures; Chris@23: } Chris@23: Chris@23: PercussionOnsetDetector::FeatureSet Chris@23: PercussionOnsetDetector::getRemainingFeatures() Chris@23: { Chris@23: return FeatureSet(); Chris@23: } Chris@23: