{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,6,17]],"date-time":"2026-06-17T09:28:43Z","timestamp":1781688523855,"version":"3.54.5"},"reference-count":53,"publisher":"Elsevier BV","issue":"1-2","license":[{"start":{"date-parts":[[2000,9,1]],"date-time":"2000-09-01T00:00:00Z","timestamp":967766400000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.elsevier.com\/tdm\/userlicense\/1.0\/"}],"content-domain":{"domain":[],"crossmark-restriction":false},"short-container-title":["Speech Communication"],"published-print":{"date-parts":[[2000,9]]},"DOI":"10.1016\/s0167-6393(00)00028-5","type":"journal-article","created":{"date-parts":[[2003,4,5]],"date-time":"2003-04-05T03:57:58Z","timestamp":1049515078000},"page":"127-154","source":"Crossref","is-referenced-by-count":222,"title":["Prosody-based automatic segmentation of speech into sentences and topics"],"prefix":"10.1016","volume":"32","author":[{"given":"Elizabeth","family":"Shriberg","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Andreas","family":"Stolcke","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Dilek","family":"Hakkani-T\u00fcr","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"G\u00f6khan","family":"T\u00fcr","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"78","reference":[{"key":"10.1016\/S0167-6393(00)00028-5_BIB1","unstructured":"Allan, J., Carbonell, J., Doddington, G., Yamron, J., Yang, Y., 1998. Topic detection and tracking pilot study: Final report. In: Proceedings DARPA Broadcast News Transcription and Understanding Workshop, Lansdowne, VA, February. Morgan Kaufmann, Los Altos, CA, pp. 194\u2013218"},{"issue":"7","key":"10.1016\/S0167-6393(00)00028-5_BIB2","doi-asserted-by":"crossref","first-page":"1001","DOI":"10.1109\/29.32278","article-title":"A tree-based statistical language model for natural language speech recognition","volume":"37","author":"Bahl","year":"1989","journal-title":"IEEE Transactions on Acoustics, Speech and Signal Processing"},{"issue":"1","key":"10.1016\/S0167-6393(00)00028-5_BIB3","doi-asserted-by":"crossref","first-page":"164","DOI":"10.1214\/aoms\/1177697196","article-title":"A maximization technique occurring in the statistical analysis of probabilistic functions in Markov chains","volume":"41","author":"Baum","year":"1970","journal-title":"The Annals of Mathematical Statistics"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB4","doi-asserted-by":"crossref","unstructured":"Beeferman, D., Berger, A., Lafferty, J., 1999. Statistical models for text segmentation. Machine Learning 34(1\u20133), 177\u2013210 (Special Issue on Natural Language Learning)","DOI":"10.1023\/A:1007506220214"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB5","series-title":"Classification and Regression Trees","author":"Breiman","year":"1984"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB6","series-title":"Questions of Intonation","author":"Brown","year":"1980"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB7","doi-asserted-by":"crossref","first-page":"274","DOI":"10.1159\/000261667","article-title":"Textual aspects of prosody in Swedish","volume":"39","author":"Bruce","year":"1982","journal-title":"Phonetica"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB8","series-title":"Introduction to IND Version 2.1 and Recursive Partitioning","author":"Buntine","year":"1992"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB9","unstructured":"Cieri, C., Graff, D., Liberman, M., Martey, N., Strassell, S., 1999. The TDT-2 text and speech corpus. In: Proceedings DARPA Broadcast News Workshop, Herndon, VA, February. Morgan Kaufmann, Los Altos, CA, pp. 57\u201360"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB10","unstructured":"Conversational Speech Recognition Workshop DARPA Hub-5E Evaluation, Baltimore, MD, May 1997"},{"issue":"2","key":"10.1016\/S0167-6393(00)00028-5_BIB11","first-page":"137","article-title":"Automatic stochastic tagging of natural language texts","volume":"21","author":"Dermatas","year":"1995","journal-title":"Computational Linguistics"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB12","doi-asserted-by":"crossref","unstructured":"Digalakis, V., Murveit, H., 1994. GENONES: An algorithm for optimizing the degree of tying in a large vocabulary hidden Markov model based speech recognizer. In: Proceedings of the IEEE Conference on Acoustics, Speech, and Signal Processing, Adelaide, Australia. Vol. 1, pp. 537\u2013540","DOI":"10.1109\/ICASSP.1994.389212"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB13","unstructured":"Doddington, G., 1998. The Topic Detection and Tracking Phase 2 (TDT2) evaluation plan. In: Proceedings DARPA Broadcast News Transcription and Understanding Workshop, Lansdowne, VA, February. Morgan Kaufmann, pp. 223\u2013229 (Revised version available from http:\/\/www.nistgov\/speech\/tdt98\/tdt98.htm.)"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB14","unstructured":"Entropic Research Laboratory, 1993. ESPS Version 5.0 Programs Manual, Washington, DC, August"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB15","doi-asserted-by":"crossref","unstructured":"Godfrey, J.J., Holliman, E.C., McDaniel. J., 1992. SWITCHBOARD: Telephone speech corpus for research and development. In: Proceedings of the IEEE Conference on Acoustics, Speech, and Signal Processing, San Francisco, March. Vol. 1, pp. 517\u2013520","DOI":"10.1109\/ICASSP.1992.225858"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB16","unstructured":"Graff, D., 1997. The 1996 Broadcast News speech and language-model corpus. In: Proceedings DARPA Speech Recognition Workshop, Chantilly, VA, February. Morgan Kaufmann, Los Altos, CA, pp. 11\u201314"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB17","doi-asserted-by":"crossref","unstructured":"Grosz, B., Hirschberg, J., 1992. Some intonational characteristics of discourse structure. In: Ohala, J.J., Nearey, T.M., Derwing, B.L., Hodge, M.M., Wiebe, G.E. (Eds.), Proceedings of the International Conference on Spoken Language Processing, Banff, Canada, October. Vol. 1, pp. 429\u2013432","DOI":"10.21437\/ICSLP.1992-103"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB18","doi-asserted-by":"crossref","unstructured":"Hakkani-T\u00fcr, D., T\u00fcr, G., Stolcke, A., Shriberg. E., 1999. Combining words and prosody for information extraction from speech. In: Proceedings of the Sixth European Conference on Speech Communication and Technology, Budapest, September. Vol. 5, pp. 1991\u20131994","DOI":"10.21437\/Eurospeech.1999-439"},{"issue":"1","key":"10.1016\/S0167-6393(00)00028-5_BIB19","first-page":"33","article-title":"TexTiling: segmenting text info multi-paragraph subtopic passages","volume":"23","author":"Hearst","year":"1997","journal-title":"Computational Linguistics"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB20","doi-asserted-by":"crossref","unstructured":"Heeman, P., Allen, J., 1997. Intonational boundaries, speech repairs, and discourse markers: modeling spoken dialog. In: Proceedings of the 35th Annual Meeting of the Association for Computational Linguistics and Eighth Conference of the European Chapter of the Association for Computational Linguistics, Madrid, July","DOI":"10.3115\/976909.979650"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB21","doi-asserted-by":"crossref","unstructured":"Hirschberg, J., Nakatani, C., 1996. A prosodic analysis of discourse segments in direction-giving monologues. In: Proceedings of the 34th Annual Meeting of the Association for Computational Linguistics, Santa Cruz, CA, June. pp. 286\u2013293","DOI":"10.3115\/981863.981901"},{"issue":"3","key":"10.1016\/S0167-6393(00)00028-5_BIB22","doi-asserted-by":"crossref","first-page":"400","DOI":"10.1109\/TASSP.1987.1165125","article-title":"Estimation of probabilities from sparse data for the language model component of a speech recognizer","volume":"35","author":"Katz","year":"1987","journal-title":"IEEE Transactions on Acoustics Speech, and Signal Processing"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB23","doi-asserted-by":"crossref","unstructured":"Koopmans-van Beinum, F.J., van Donzel, M.E., 1996. Relationship between discourse structure and dynamic speech rate. In: Bunnell, H.T., Idsardi, W. (Eds.), Proceedings of the International Conference on Spoken Language Processing, Vol. 3, Philadelphia, October, pp. 1724\u20131727","DOI":"10.1109\/ICSLP.1996.607960"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB24","doi-asserted-by":"crossref","unstructured":"Kozima, H., 1993. Text segmentation based on similarity between words. In: Proceedings of the 31st Annual Meeting of the Association for Computational Linguistics, Ohio State University, Columbus, Ohio, June. pp. 286\u2013288","DOI":"10.3115\/981574.981616"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB25","unstructured":"Kubala, F., Schwartz, R., Stone, R., Weischedel, R., 1998. Named entity extraction from speech. In: Proceedings DARPA Broadcast News Transcription and Understanding Workshop, Lansdowne, VA, February. Morgan Kaufmann, Los Altos, CA, pp. 287\u2013292"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB26","unstructured":"Lehiste, I., 1979. Perception of sentence and paragraph boundaries. In: Lindblom, B., \u00d6hman, S. (Eds.), Frontiers of Speech Communication Research. Academic, London, pp. 191\u2013201"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB27","doi-asserted-by":"crossref","unstructured":"Lehiste, I., 1980. The phonetic structure of paragraphs. In: Nooteboom, S., Cohen, A. (Eds.), Structure and Process in Speech Perception. Springer, Berlin, pp. 195\u2013206","DOI":"10.1007\/978-3-642-81000-8_12"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB28","doi-asserted-by":"crossref","unstructured":"Liu, D., Kubala, F., 1999. Fast speaker change detection for Broadcast News transcription and indexing. In: Proceedings of the Sixth European Conference on Speech Communication and Technology, Budapest, September. Vol. 3, pp. 1031\u20131034","DOI":"10.21437\/Eurospeech.1999-167"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB29","unstructured":"LVCSR Hub-5 Workshop, Linthicum Heights, MD, June 1999"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB30","unstructured":"Meteer, M., Taylor, A., MacIntyre, R., Iyer, R., 1995. Dysfluency annotation stylebook for the Switchboard corpus. Distributed by LDC, ftp:\/\/ftp.cis.upenn.edu\/pub\/treebank\/swbd\/doc\/DFL-book.ps, February (Revised June 1995 by Ann Taylor)"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB31","doi-asserted-by":"crossref","unstructured":"Nakajima, S., Tsukada, H., 1997. Prosodic features of utterances in task-oriented dialogues. In: Sagisaka, Y., Campbell, N., Higuchi, N. (Eds.), Computing Prosody: Computational Models for Processing Spontaneous Speech. Springer, New York, Chapter 7, pp. 81\u201394","DOI":"10.1007\/978-1-4612-2258-3_7"},{"issue":"2","key":"10.1016\/S0167-6393(00)00028-5_BIB32","first-page":"241","article-title":"Adaptive multilingual sentence boundary disambiguation","volume":"23","author":"Palmer","year":"1997","journal-title":"Computational Linguistics"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB33","doi-asserted-by":"crossref","unstructured":"Przybocki, M.A., Martin, A.F., 1999. The 1999 NIST speaker recognition evaluation, using summed two-channel telephone data for speaker detection and speaker tracking. In: Proceedings of the Sixth European Conference on Speech Communication and Technology, Budapest. Vol. 5, pp. 2215\u20132218","DOI":"10.21437\/Eurospeech.1999-491"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB34","unstructured":"Sankar, A., Weng, F., Rivlin, Z., Stolcke, A., Gadde, R.R., 1998. The development of SRI's 1997 Broadcast News transcription system. In: Proceedings DARPA Broadcast News Transcription and Understanding Workshop, Lansdowne, VA, February. Morgan Kaufmann, Los Altos, CA, pp. 91\u201396"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB35","unstructured":"Shriberg, E., 1999. Phonetic consequences of speech disfluency. In: Proceedings of the XIVth International Congress on Phonetic Sciences, San Francisco. pp. 619\u2013622"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB36","doi-asserted-by":"crossref","unstructured":"Shriberg, E., Bates, R., Stolcke, A., 1997. A prosody-only decision-tree model for disfluency detection. In: Kokkinakis, G. Fakotakis, N., Dermatas, E. (Eds.), Proceedings of the Fifth European Conference on Speech Communication and Technology, Rhodes, Greece, September. Vol. 5, pp. 2383\u20132386","DOI":"10.21437\/Eurospeech.1997-626"},{"issue":"3\u20134","key":"10.1016\/S0167-6393(00)00028-5_BIB37","first-page":"439","article-title":"Can prosody aid the automatic classification of dialog acts in conversational speech?","volume":"41","author":"Shriberg","year":"1998","journal-title":"Language and Speech"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB38","unstructured":"Silverman, K., 1987. The structure and processing of fundamental frequency contours. Ph.D. thesis, Cambridge University, Cambridge, UK"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB39","doi-asserted-by":"crossref","first-page":"180","DOI":"10.1159\/000261938","article-title":"Beyond sentence prosody: paragraph intonation in Dutch","volume":"50","author":"Sluijter","year":"1994","journal-title":"Phonetica"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB40","doi-asserted-by":"crossref","unstructured":"S\u00f6nmez, K., Shriberg, E., Heck, L., Weintraub, M., 1998. Modeling dynamic prosodic variation for speaker verification. In: Mannell, R.H., Robert-Ribes, J. (Eds.), Proceedings of the International Conference on Spoken Language Processing, Sydney, December. Australian Speech Science and Technology Association, Vol. 7, pp. 3189\u20133192","DOI":"10.21437\/ICSLP.1998-254"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB41","doi-asserted-by":"crossref","unstructured":"S\u00f6nmez, K., Heck, L., Weintraub, M., 1999. Speaker tracking and detection with multiple speakers. In: Proceedings of the Sixth European Conference on Speech Communication and Technology, Budapest. Vol. 5, pp. 2219\u20132222","DOI":"10.21437\/Eurospeech.1999-492"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB42","doi-asserted-by":"crossref","unstructured":"Stolcke, A., Shriberg, E., 1996. Automatic linguistic segmentation of conversational speech. In: Bunnell, H.T., Idsardi, W. (Eds.), Proceedings of the International Conference on Spoken Language Processing, Philadelphia. Vol. 2, pp. 1005\u20131008","DOI":"10.1109\/ICSLP.1996.607773"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB43","doi-asserted-by":"crossref","unstructured":"Stolcke, A., Shriberg, E., Bates, R., Ostendorf, M., Hakkani, D., Plauch\u00e9, M., T\u00fcr, G., Lu, Y., 1998. Automatic detection of sentence boundaries and disfluencies based on recognized words. In: Mannell, R.H., Robert-Ribes, J. (Eds.), Proceedings of the International Conference on Spoken Language Processing, Sydney. Australian Speech Science and Technology Association, Vol. 5, pp. 2247\u20132250","DOI":"10.21437\/ICSLP.1998-486"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB44","unstructured":"Stolcke, A., Shriberg, E., Hakkani-T\u00fcr, D., T\u00fcr, G., Rivlin, Z., S\u00f6nmez, K., 1999. Combining words and speech prosody for automatic topic segmentation. In: Proceedings DARPA Broadcast News Workshop, Herndon, VA, February. Morgan Kaufmann, Los Altos, CA, pp. 61\u201364"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB45","doi-asserted-by":"crossref","first-page":"514","DOI":"10.1121\/1.418114","article-title":"Prosodic features at discourse boundaries of different strength","volume":"101","author":"Swerts","year":"1997","journal-title":"Journal of the Acoustical Society of America"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB46","doi-asserted-by":"crossref","first-page":"21","DOI":"10.1177\/002383099403700102","article-title":"Prosody as a marker of information flow in spoken discourse","volume":"37","author":"Swerts","year":"1994","journal-title":"Language and Speech"},{"issue":"1","key":"10.1016\/S0167-6393(00)00028-5_BIB47","doi-asserted-by":"crossref","first-page":"25","DOI":"10.1016\/S0167-6393(97)00011-3","article-title":"Prosodic and lexical indications of discourse structure in human\u2013machine interactions","volume":"22","author":"Swerts","year":"1997","journal-title":"Speech Communication"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB48","unstructured":"Talkin, D., 1995. A robust algorithm for pitch tracking (RAPT). In: Klein, W.B., Paliwal, K.K. (Eds.), Speech Coding and Synthesis. Elsevier, New York"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB49","doi-asserted-by":"crossref","first-page":"1205","DOI":"10.1121\/1.392187","article-title":"Intonation and text in Standard Dutch","volume":"77","author":"Thorsen","year":"1985","journal-title":"Journal of the Acoustical Society of America"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB50","doi-asserted-by":"crossref","unstructured":"T\u00fcr, G., Hakkani-T\u00fcr, D., Stolcke, A., Shriberg, E., To appear. Integrating prosodic and lexical cues for automatic topic segmentation, Computational Linguistics","DOI":"10.1162\/089120101300346796"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB51","doi-asserted-by":"crossref","unstructured":"Vaissi\u00e8re. J., 1983. Language-independent prosodic features. In Cutler, A., Ladd, D.R. (Eds.), Prosody: Models and Measurements. Springer, Berlin, Chapter 5, pp. 53\u201366","DOI":"10.1007\/978-3-642-69103-4_5"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB52","doi-asserted-by":"crossref","first-page":"260","DOI":"10.1109\/TIT.1967.1054010","article-title":"Error bounds for convolutional codes and an asymptotically optimum decoding algorithm","volume":"13","author":"Viterbi","year":"1967","journal-title":"IEEE Transactions on Information Theory"},{"key":"10.1016\/S0167-6393(00)00028-5_BIB53","doi-asserted-by":"crossref","unstructured":"Yamron, J.P., Carp, I., Gillick, L., Lowe, S., van Mulbregt. P., 1998. A hidden Markov model approach to text segmentation and event tracking. In: Proceedings of the IEEE Conference on Acoustics, Speech, and Signal Processing, Seattle, WA. Vol. 1, pp. 333\u2013336","DOI":"10.1109\/ICASSP.1998.674435"}],"container-title":["Speech Communication"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/api.elsevier.com\/content\/article\/PII:S0167639300000285?httpAccept=text\/xml","content-type":"text\/xml","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/api.elsevier.com\/content\/article\/PII:S0167639300000285?httpAccept=text\/plain","content-type":"text\/plain","content-version":"vor","intended-application":"text-mining"}],"deposited":{"date-parts":[[2023,4,9]],"date-time":"2023-04-09T02:52:37Z","timestamp":1681008757000},"score":1,"resource":{"primary":{"URL":"https:\/\/linkinghub.elsevier.com\/retrieve\/pii\/S0167639300000285"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2000,9]]},"references-count":53,"journal-issue":{"issue":"1-2","published-print":{"date-parts":[[2000,9]]}},"alternative-id":["S0167639300000285"],"URL":"https:\/\/doi.org\/10.1016\/s0167-6393(00)00028-5","relation":{},"ISSN":["0167-6393"],"issn-type":[{"value":"0167-6393","type":"print"}],"subject":[],"published":{"date-parts":[[2000,9]]}}}