{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2024,10,22]],"date-time":"2024-10-22T18:37:23Z","timestamp":1729622243560,"version":"3.28.0"},"reference-count":40,"publisher":"IEEE","content-domain":{"domain":[],"crossmark-restriction":false},"short-container-title":[],"published-print":{"date-parts":[[2015,12]]},"DOI":"10.1109\/asru.2015.7404840","type":"proceedings-article","created":{"date-parts":[[2016,2,12]],"date-time":"2016-02-12T08:55:42Z","timestamp":1455267342000},"page":"525-532","source":"Crossref","is-referenced-by-count":6,"title":["Improving robustness against reverberation for automatic speech recognition"],"prefix":"10.1109","author":[{"given":"Vikramjit","family":"Mitra","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Julien","family":"Van Hout","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Wen","family":"Wang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Martin","family":"Graciarena","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Mitchell","family":"McLaren","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Horacio","family":"Franco","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Dimitra","family":"Vergyri","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"263","reference":[{"key":"ref39","doi-asserted-by":"publisher","DOI":"10.1109\/ASRU.2015.7404833"},{"key":"ref38","article-title":"The SRI March 2000 Hub-5 Conversational Speech Transcription System","author":"stolcke","year":"2000","journal-title":"Proc NIST Speech Transcription Workshop"},{"key":"ref33","first-page":"901","article-title":"SRILM&#x2014;An Extensible Language Modeling Toolkit","author":"stolcke","year":"2002","journal-title":"Proc of ICSLP 2002"},{"key":"ref32","article-title":"Time Frequency Convolution Nets for Robust Speech Recognition","author":"mitra","year":"2015","journal-title":"Submitted to ASRU"},{"key":"ref31","doi-asserted-by":"publisher","DOI":"10.1109\/89.841209"},{"key":"ref30","doi-asserted-by":"publisher","DOI":"10.1109\/TASSP.1980.1163453"},{"key":"ref37","doi-asserted-by":"publisher","DOI":"10.1109\/ASRU.1997.659110"},{"key":"ref36","doi-asserted-by":"publisher","DOI":"10.1109\/SLT.2014.7078633"},{"key":"ref35","article-title":"RNNLM&#x2014;Recurrent Neural Network Language Modeling Toolkit","author":"mikolov","year":"2011","journal-title":"Proc of ASRU"},{"key":"ref34","doi-asserted-by":"publisher","DOI":"10.1109\/ICASSP.2009.4960587"},{"key":"ref10","doi-asserted-by":"publisher","DOI":"10.1109\/TASL.2008.2010214"},{"key":"ref40","article-title":"The kaldi speech recognition toolkit","author":"povey","year":"2011","journal-title":"Proc ASRU"},{"key":"ref11","article-title":"Use Of Multiple Front-Ends and I-Vector-Based Speaker Adaptation for Robust Speech Recognition","author":"alam","year":"2014","journal-title":"Proc of REVERB Challenge"},{"key":"ref12","article-title":"Linear Prediction-Based Dereverberation with Advanced Speech Enhancement and Recognition Technologies for the REVERB Challenge","author":"delcroix","year":"2014","journal-title":"Proc of REVERB Challenge"},{"key":"ref13","article-title":"Robust Features and System Fusion for Reverberation-Robust Speech Recognition","author":"mitra","year":"2014","journal-title":"Proc of REVERB Challenge"},{"key":"ref14","article-title":"Combating Reverberation in Large Vocabulary Continuous Speech Recognition","author":"mitra","year":"2015","journal-title":"Proc of Interspeech"},{"key":"ref15","article-title":"Vocal Tract Length Normalization for LVCSR","author":"zhan","year":"1997","journal-title":"Tech Rep CMU-LTI-97-150"},{"key":"ref16","article-title":"Evaluating Robust Features on Deep Neural Networks for Speech Recognition in Noisy and Channel Mismatched Conditions","author":"mitra","year":"2014","journal-title":"Proc of Interspeech"},{"key":"ref17","doi-asserted-by":"publisher","DOI":"10.1109\/ICASSP.2012.6288864"},{"key":"ref18","first-page":"1727","article-title":"Recent Improvements in SRI's Keyword Detection System for Noisy Audio","author":"van-hout","year":"2014","journal-title":"Proc of Interspeech"},{"key":"ref19","doi-asserted-by":"publisher","DOI":"10.1109\/ICASSP.2013.6639347"},{"key":"ref28","doi-asserted-by":"publisher","DOI":"10.1109\/ICASSP.2012.6288824"},{"key":"ref4","doi-asserted-by":"publisher","DOI":"10.1109\/WASPAA.2013.6701894"},{"key":"ref27","doi-asserted-by":"publisher","DOI":"10.1121\/1.1396325"},{"key":"ref3","doi-asserted-by":"publisher","DOI":"10.1109\/ICASSP.2013.6639100"},{"key":"ref6","doi-asserted-by":"crossref","first-page":"429","DOI":"10.1007\/BF02999431","article-title":"Combined Acoustic Echo Cancellation, Dereverberation and Noise Reduction: A Two Microphone Approach","volume":"49","author":"martin","year":"1994","journal-title":"Annales of Telecommunications"},{"key":"ref29","doi-asserted-by":"publisher","DOI":"10.1109\/ICASSP.2014.6853898"},{"key":"ref5","doi-asserted-by":"crossref","DOI":"10.1007\/978-3-662-04619-7","author":"brandstein","year":"2001","journal-title":"Microphone Arrays Signal Processing Techniques and Applications"},{"key":"ref8","doi-asserted-by":"publisher","DOI":"10.1109\/TSA.2005.858066"},{"key":"ref7","first-page":"1","article-title":"Single Channel Blind Dereverberation Based on Auto-Correlation Functions of Frame-Wise Time Sequences of Frequency Components","author":"ohta","year":"2006","journal-title":"Proc of IWAENC"},{"key":"ref2","doi-asserted-by":"crossref","DOI":"10.21437\/Interspeech.2011-169","article-title":"Conversational Speech Transcription Using Context-Dependent Deep Neural Networks","author":"seide","year":"2011","journal-title":"Proc of Interspeech"},{"key":"ref9","doi-asserted-by":"publisher","DOI":"10.1109\/ICASSP.2007.366926"},{"key":"ref1","doi-asserted-by":"publisher","DOI":"10.1109\/TASL.2011.2109382"},{"key":"ref20","doi-asserted-by":"publisher","DOI":"10.1109\/ASRU.2013.6707705"},{"key":"ref22","article-title":"The Automatic Speech Recognition in Reverberant Environments (ASpIRE) Challenge","author":"harper","year":"2015","journal-title":"Proc of ASRU"},{"key":"ref21","doi-asserted-by":"publisher","DOI":"10.1109\/TASL.2010.2064307"},{"key":"ref24","doi-asserted-by":"publisher","DOI":"10.1109\/ICASSP.2011.5947380"},{"key":"ref23","first-page":"709","article-title":"All for One: Feature Combination for Highly Channel-Degraded Speech Activity Detection","author":"graciarena","year":"2013","journal-title":"Proc of Interspeech"},{"key":"ref26","doi-asserted-by":"publisher","DOI":"10.1121\/1.409836"},{"key":"ref25","first-page":"886","article-title":"Damped Oscillator Cepstral Coefficients for Robust Speech Recognition","author":"mitra","year":"2013","journal-title":"Proc of Interspeech"}],"event":{"name":"2015 IEEE Workshop on Automatic Speech Recognition and Understanding (ASRU)","start":{"date-parts":[[2015,12,13]]},"location":"Scottsdale, AZ, USA","end":{"date-parts":[[2015,12,17]]}},"container-title":["2015 IEEE Workshop on Automatic Speech Recognition and Understanding (ASRU)"],"original-title":[],"link":[{"URL":"http:\/\/xplorestaging.ieee.org\/ielx7\/7397480\/7404758\/07404840.pdf?arnumber=7404840","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2022,6,3]],"date-time":"2022-06-03T21:24:39Z","timestamp":1654291479000},"score":1,"resource":{"primary":{"URL":"http:\/\/ieeexplore.ieee.org\/document\/7404840\/"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2015,12]]},"references-count":40,"URL":"https:\/\/doi.org\/10.1109\/asru.2015.7404840","relation":{},"subject":[],"published":{"date-parts":[[2015,12]]}}}