{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2025,5,3]],"date-time":"2025-05-03T04:02:18Z","timestamp":1746244938723,"version":"3.40.4"},"reference-count":33,"publisher":"Springer Science and Business Media LLC","issue":"3","license":[{"start":{"date-parts":[[2014,3,21]],"date-time":"2014-03-21T00:00:00Z","timestamp":1395360000000},"content-version":"tdm","delay-in-days":0,"URL":"http:\/\/www.springer.com\/tdm"}],"content-domain":{"domain":[],"crossmark-restriction":false},"short-container-title":["Int J Speech Technol"],"published-print":{"date-parts":[[2014,9]]},"DOI":"10.1007\/s10772-014-9229-5","type":"journal-article","created":{"date-parts":[[2014,3,20]],"date-time":"2014-03-20T10:15:47Z","timestamp":1395310547000},"page":"259-269","source":"Crossref","is-referenced-by-count":0,"title":["Segmentation, indexing and retrieval of TV broadcast news bulletins using Gaussian mixture models and vector quantization codebooks"],"prefix":"10.1007","volume":"17","author":[{"given":"K. Sreenivasa","family":"Rao","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Ketan","family":"Pachpande","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2014,3,21]]},"reference":[{"key":"9229_CR1","doi-asserted-by":"crossref","unstructured":"Antonelli, M., Rizzi, A., & del Vescovo, G. (2010, Dec). A query by humming system for music information retrieval. In: Intelligent Systems Design and Applications (ISDA), 10th International Conference (pp.586\u2013591).","DOI":"10.1109\/ISDA.2010.5687200"},{"key":"9229_CR2","doi-asserted-by":"crossref","unstructured":"Bengherabi, M., & Sehad, A. (Apr. 2006). Development and evaluation of automatic-speaker based-audio identification and segmentation for broadcast news recordings indexation. Information and Communication Technologies, 1, 1230\u20131235.","DOI":"10.1109\/ICTTA.2006.1684553"},{"key":"9229_CR3","doi-asserted-by":"crossref","unstructured":"Butko, T., & Nadeu, C. (2011, May). Audio segmentation of broadcast news: A hierarchical system with feature selection for the Albayzin-2010 evaluation. In: Proceedings of IEEE International Conference on Acoustics, Speech and Signal Processing (pp.357\u2013360)","DOI":"10.1109\/ICASSP.2011.5946414"},{"key":"9229_CR4","unstructured":"Chen, S., & Gopalakrishnan, P. S. (1998). Speaker, environment and channel change detection and clustering via the bayesian information criterion. In Proceedings of DARPA Broadcast News Transcription and Understanding Workshop."},{"issue":"12","key":"9229_CR5","doi-asserted-by":"crossref","first-page":"111","DOI":"10.1016\/S0167-6393(00)00027-3","volume":"32","author":"P Delacourt","year":"2000","unstructured":"Delacourt, P., & Wellekens, C. (2000). Distbic: A speaker-based segmentation for audio data indexing. Speech Communication, 32(12), 111\u2013126.","journal-title":"Speech Communication"},{"key":"9229_CR6","doi-asserted-by":"crossref","unstructured":"Dhananjaya, N., Prasad, S. G., and Yegnanarayana, B. (2004, Nov). Speaker segmentation based on subsegmental features and neural network models, 11th International Conference on Neural Information Processing (ICONIP-2004), vol. 50, (pp.1210\u20131215).","DOI":"10.1007\/978-3-540-30499-9_188"},{"key":"9229_CR7","doi-asserted-by":"crossref","first-page":"153","DOI":"10.1016\/j.specom.2007.08.003","volume":"50","author":"N Dhananjaya","year":"2008","unstructured":"Dhananjaya, N., & Yegnanarayana, B. (2008). Speaker change detection in casual conversations using excitation source features. Speech Communication, 50, 153\u2013161.","journal-title":"Speech Communication"},{"key":"9229_CR8","doi-asserted-by":"crossref","unstructured":"Foote, J. (2000). Automatic audio segmentation using a measure of audio novelty. In: Proceedings of International Conference on Multimedia and Expo, textit1, (pp.452\u2013455)","DOI":"10.1109\/ICME.2000.869637"},{"key":"9229_CR9","doi-asserted-by":"crossref","unstructured":"Gish, H., Siu, M.-H., & Rohlicek, R. (1991). Segregation of speakers for speech recognition and speaker identification. In Proceedings of IEEE International Conference acoust, speech and signal processing, 2,(pp.873\u2013876).","DOI":"10.1109\/ICASSP.1991.150477"},{"key":"9229_CR10","doi-asserted-by":"crossref","unstructured":"Hauptmann, A.G., and Witbrock, M. J. (1998, April). Story segmentation and detection of commercials in broadcast news video. In Proceedings of IEEE International Forum Research and Technology Advances in Digital Libraries, Santa Barbara, CA, USA (pp.168\u2013179)","DOI":"10.1109\/ADL.1998.670392"},{"issue":"654\u2013655","key":"9229_CR11","first-page":"29","volume":"46","author":"Q-H He","year":"2010","unstructured":"He, Q.-H., Yang, J.-C., Li, Y.-X., He, J., Zhang, X.-Y., & Li, W. (2010). Combining GMM, Jensen\u2019s inequality and BIC for speaker indexing. Electronics Letters, 46(654\u2013655), 29.","journal-title":"Electronics Letters"},{"key":"9229_CR12","doi-asserted-by":"crossref","unstructured":"Huang, R., & Hansen, J. H. L. (2004, May). Advances in unsupervised audio segmentation for the broadcast news and ngsw corpora. In: Proceedings of IEEE International Conference on Acoustics, Speech and Signal Processing 1, (pp.741\u2013744).","DOI":"10.1109\/ICASSP.2004.1326092"},{"key":"9229_CR13","doi-asserted-by":"crossref","unstructured":"Karydis, A. P.: I Nanopoulos and Y. Manolopoulos. (2005, Jan). Audio indexing for efficient music information retrieval. In: Multimedia Modelling Conference, (pp.22\u201329)","DOI":"10.1109\/MMMC.2005.22"},{"key":"9229_CR14","doi-asserted-by":"crossref","unstructured":"Kemp, T., Schmidt, M., Westphal, M., & Waibel, A. (2000). Strategies for automatic segmentation of audio data. In Proceedings of IEEE International Conference Acoust Speech Signal Processing, 3, 1423\u20131426.","DOI":"10.1109\/ICASSP.2000.861862"},{"key":"9229_CR15","doi-asserted-by":"crossref","first-page":"920","DOI":"10.1109\/TASL.2008.925152","volume":"16","author":"M Kotti","year":"2008","unstructured":"Kotti, M., Benetos, E., & Kotropoulos, C. (2008). Computationally efficient and robust bic-based speaker segmentation. IEEE Transactions on Audio, Speech and Language Processing, 16, 920\u2013933.","journal-title":"IEEE Transactions on Audio, Speech and Language Processing"},{"key":"9229_CR16","unstructured":"Lei, W.: Unsupervised techniques for audio content analysis and summarization. PhD thesis, School of Computer Engineering, Nanyang Technological University, Singapore, May 2008."},{"key":"9229_CR17","doi-asserted-by":"crossref","unstructured":"Li, D., Sethi, I. K., Dimitrova, N., & McGee, T. (Apr. 2001). Classification of general audio data for content-based retrieval. Pattern Recognition Letters, 22, 533\u2013544.","DOI":"10.1016\/S0167-8655(00)00119-7"},{"key":"9229_CR18","doi-asserted-by":"crossref","unstructured":"Lu, L, Jiang, H., & Zhang, H. J. (2001, Oct). A robust audio classification and segmentation method. In: Proceedings of the ninth ACM International Conference in Multimedia, Ottawa, Canada (pp.203\u2013211).","DOI":"10.1145\/500141.500173"},{"key":"9229_CR19","doi-asserted-by":"crossref","first-page":"269","DOI":"10.1023\/A:1012491016871","volume":"15","author":"G Lu","year":"2001","unstructured":"Lu, G. (2001). Indexing and retrieval of audio: A survey. Multimedia Tools and Applications, 15, 269\u2013290.","journal-title":"Multimedia Tools and Applications"},{"key":"9229_CR20","doi-asserted-by":"crossref","first-page":"1338","DOI":"10.1109\/5.880087","volume":"88","author":"J Makhoul","year":"2000","unstructured":"Makhoul, J., Kubala, F., Leek, T., Leu, D., Nguyen, L., Schwartz, R., et al. (2000). Speech and language technologies for audio indexing and retrieval. Processing of the IEEE, 88, 1338\u20131353.","journal-title":"Processing of the IEEE"},{"key":"9229_CR21","doi-asserted-by":"crossref","unstructured":"Meinedo, H., & Neto, J. (2003, April). Audio segmentation, classification and clustering in a broadcast news task. In: Proceedings of IEEE International Conference on Acoustics Speech and Signal Processing, 2, (pp.5\u20138.)","DOI":"10.1109\/ICASSP.2003.1202280"},{"key":"9229_CR22","doi-asserted-by":"crossref","unstructured":"Nwe, T. L., & Li, H. (2005, Mar). Broadcast news segmentation by audio type analysis. In: Proceedings of IEEE International Conference on Acoustics Speech snd Signal Processing 2, (pp.1065\u20131068).","DOI":"10.1109\/ICASSP.2005.1415592"},{"key":"9229_CR23","doi-asserted-by":"crossref","first-page":"69","DOI":"10.1109\/MSP.2006.1621450","volume":"23","author":"K Ohtsuki","year":"2006","unstructured":"Ohtsuki, K., Bessho, K., Matsuo, Y., Matsunaga, S., & Hayashi, Y. (2006). Automatic multimedia indexing: combining audio, speech, and visual information to index broadcast news. Signal Processing Magazine IEEE, 23, 69\u201378.","journal-title":"Signal Processing Magazine IEEE"},{"key":"9229_CR24","doi-asserted-by":"crossref","unstructured":"Perez-Freire, L., & Garcia-Mateo, C. (2004, May). A multimedia approach for audio segmentation in TV broadcast news. In: Proceedings of IEEE International Conference on Acoustics, Speech and Signal Processing, vol. 1.","DOI":"10.1109\/ICASSP.2004.1325999"},{"key":"9229_CR25","doi-asserted-by":"crossref","unstructured":"Rao, K. S., Pachpande, K., Reddy, V. R., & Maity, S. (2012, Feb). Segmentation of tv broadcast news using speaker specific information. In: National Conference on Communications (NCC-2012), IIT Kharagpur, Kharagpur, India.","DOI":"10.1109\/NCC.2012.6176848"},{"key":"9229_CR26","unstructured":"Reiss, J., Aucouturier, J. J., & Sandler, M. (2001). Efficient multidimensional searching routines for music information retrieval. International Society of Musical, Information Retrieval, pp.163\u2013171, 2001."},{"key":"9229_CR27","doi-asserted-by":"crossref","first-page":"5","DOI":"10.1016\/S0167-6393(00)00020-0","volume":"32","author":"S Renals","year":"2000","unstructured":"Renals, S., Abberley, D., Kirby, D., & Robinson, T. (2000). Indexing and retrieval of broadcast news. Speech Communication, 32, 5\u201320.","journal-title":"Speech Communication"},{"issue":"6","key":"9229_CR28","doi-asserted-by":"crossref","first-page":"1894","DOI":"10.1109\/TASL.2012.2191284","volume":"20","author":"AK Vuppala","year":"2012","unstructured":"Vuppala, A. K., Yadav, J., Chakrabarti, S., & Rao, K. S. (2012). Vowel onset point detection for low bit rate coded speech. IEEE Transactions on Audio, Speech and Language Processing, 20(6), 1894\u20131903.","journal-title":"IEEE Transactions on Audio, Speech and Language Processing"},{"issue":"2","key":"9229_CR29","doi-asserted-by":"crossref","first-page":"229","DOI":"10.1007\/s10772-012-9179-8","volume":"16","author":"AK Vuppala","year":"2013","unstructured":"Vuppala, A. K., Rao, K. S., & Chakrabarti, S. (2013). Vowel onset point detection for noisy speech using spectral energy at formant frequencies. International Journal of Speech Technology (Springer), 16(2), 229\u2013235.","journal-title":"International Journal of Speech Technology (Springer)"},{"key":"9229_CR30","doi-asserted-by":"crossref","unstructured":"Wu, C.-H., & Hsieh, C.-H. (Mar. 2006). Multiple change-point audio segmentation and classification using an MDL-based Gaussian model. IEEE Transactions on Audio, Speech, and Language Processing, 14, 647\u2013657.","DOI":"10.1109\/TSA.2005.852988"},{"key":"9229_CR31","doi-asserted-by":"crossref","unstructured":"Xue, H., Li, H., Gao, C., & Shi, Z. (2010). Computationally efficient audio segmentation through a multi-stage bic approach. In 3rd International Congress on Image and Signal Processing (CISP), 8, (pp.3774\u20133777).","DOI":"10.1109\/CISP.2010.5646687"},{"key":"9229_CR32","doi-asserted-by":"crossref","first-page":"441","DOI":"10.1109\/89.917689","volume":"9","author":"T Zhang","year":"2001","unstructured":"Zhang, T., & Kuo, C.-C. (2001). Audio content analysis for online audiovisual data segmentation and classification. IEEE Transactions on Speech and Audio Processing, 9, 441\u2013457.","journal-title":"IEEE Transactions on Speech and Audio Processing"},{"issue":"6","key":"9229_CR33","doi-asserted-by":"crossref","first-page":"582","DOI":"10.1007\/BF02943243","volume":"16","author":"F Zheng","year":"2001","unstructured":"Zheng, F., Zhang, G., & Song, Z. (2001). Comparison of different implementations of MFCC. Journal of Computer Science and Technology, 16(6), 582\u2013589.","journal-title":"Journal of Computer Science and Technology"}],"container-title":["International Journal of Speech Technology"],"original-title":[],"language":"en","link":[{"URL":"http:\/\/link.springer.com\/content\/pdf\/10.1007\/s10772-014-9229-5.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"http:\/\/link.springer.com\/article\/10.1007\/s10772-014-9229-5\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"http:\/\/link.springer.com\/content\/pdf\/10.1007\/s10772-014-9229-5","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,5,2]],"date-time":"2025-05-02T03:27:55Z","timestamp":1746156475000},"score":1,"resource":{"primary":{"URL":"http:\/\/link.springer.com\/10.1007\/s10772-014-9229-5"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2014,3,21]]},"references-count":33,"journal-issue":{"issue":"3","published-print":{"date-parts":[[2014,9]]}},"alternative-id":["9229"],"URL":"https:\/\/doi.org\/10.1007\/s10772-014-9229-5","relation":{},"ISSN":["1381-2416","1572-8110"],"issn-type":[{"type":"print","value":"1381-2416"},{"type":"electronic","value":"1572-8110"}],"subject":[],"published":{"date-parts":[[2014,3,21]]}}}