{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2025,4,10]],"date-time":"2025-04-10T05:37:40Z","timestamp":1744263460117,"version":"3.30.2"},"reference-count":25,"publisher":"Elsevier BV","issue":"4","license":[{"start":{"date-parts":[[2003,10,1]],"date-time":"2003-10-01T00:00:00Z","timestamp":1064966400000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.elsevier.com\/tdm\/userlicense\/1.0\/"}],"content-domain":{"domain":[],"crossmark-restriction":false},"short-container-title":["Computer Speech &amp; Language"],"published-print":{"date-parts":[[2003,10]]},"DOI":"10.1016\/s0885-2308(03)00008-1","type":"journal-article","created":{"date-parts":[[2003,4,7]],"date-time":"2003-04-07T17:19:33Z","timestamp":1049735973000},"page":"357-379","source":"Crossref","is-referenced-by-count":13,"title":["Modeling partial pronunciation variations for spontaneous Mandarin speech recognition"],"prefix":"10.1016","volume":"17","author":[{"given":"Yi","family":"Liu","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Pascale","family":"Fung","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"78","reference":[{"key":"10.1016\/S0885-2308(03)00008-1_BIB1","doi-asserted-by":"crossref","unstructured":"Byrne, W., Finke, M., Khudanpur, S., Mcdonough, J., Nock, H., Riley, M., Saraclar, M., Wooters, C. Zavaliagkos, G., 1998. Pronunciation modelling using a hand-labelled corpus for conversational speech recognition. In: Proceedings of the IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP), Settle, USA, pp. 313\u2013316","DOI":"10.1109\/ICASSP.1998.674430"},{"key":"10.1016\/S0885-2308(03)00008-1_BIB2","doi-asserted-by":"crossref","unstructured":"Byrne, W., Venkataramani, V., Kamm, T., Zheng, F., Fung, P., Liu, Y., Ruhi, U., 2001. Automatic generation of pronunciation lexicons for Mandarin spontaneous speech. In: Proceedings of the IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP), Salt Lake City, USA","DOI":"10.1109\/ICASSP.2001.940895"},{"key":"10.1016\/S0885-2308(03)00008-1_BIB3","doi-asserted-by":"crossref","unstructured":"Chen, F., 1990. Identification of contextual factors for pronunciation networks. In: Proceedings of the IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP), pp. 753\u2013756","DOI":"10.1109\/ICASSP.1990.115902"},{"key":"10.1016\/S0885-2308(03)00008-1_BIB4","doi-asserted-by":"crossref","unstructured":"Eide, E., 1999. Automatic modeling of pronunciation variations. In: Proceedings of the European Conference on Speech Communication and Technology (Eurospeech), Budapest, Hungary, pp. 451\u2013454","DOI":"10.21437\/Eurospeech.1999-116"},{"key":"10.1016\/S0885-2308(03)00008-1_BIB5","doi-asserted-by":"crossref","unstructured":"Finke, M., Fritsch, J., Koll, D., Waibel, A., 1999. Modeling and efficient decoding of large vocabulary conversational speech. In: Proceedings of the European Conference on Speech Communication and Technology (Eurospeech), Budapest, Hungary, pp. 467\u2013470","DOI":"10.21437\/Eurospeech.1999-120"},{"key":"10.1016\/S0885-2308(03)00008-1_BIB6","doi-asserted-by":"crossref","unstructured":"Fosler-Lussier, E., 1999. Dynamic pronunciation models for automatic speech recognition. Ph.D. Thesis. International Computer Science Institute, Berkeley, CA","DOI":"10.1016\/S0167-6393(99)00035-7"},{"key":"10.1016\/S0885-2308(03)00008-1_BIB7","doi-asserted-by":"crossref","unstructured":"Fokada, T., Sagisaka, Y., 1997. Automatic generation of a pronunciation dictionary based on a pronunciation network. In: Proceedings of the European Conference on Speech Communication and Technology (Eurospeech), Rhodes, Greece, pp. 2471\u20132474","DOI":"10.21437\/Eurospeech.1997-642"},{"key":"10.1016\/S0885-2308(03)00008-1_BIB8","unstructured":"Fung, P., Byrne, W., Zheng, F., Kamm, T., Liu, Y., Song, Z., Venkataramani, V., Ruhi, U., 2000. Pronunciation modeling of Mandarin casual speech. Final Report, The Johns Hopkins University Summer Workshop"},{"key":"10.1016\/S0885-2308(03)00008-1_BIB9","doi-asserted-by":"crossref","unstructured":"Hain, T., Woodland, P.C., 1999. Dynamic HMM selection for continuous speech recognition. In: Proceedings of the European Conference on Speech Communication and Technology (Eurospeech), Budapest, Hungary, pp. 1327\u20131330","DOI":"10.21437\/Eurospeech.1999-339x"},{"key":"10.1016\/S0885-2308(03)00008-1_BIB10","doi-asserted-by":"crossref","first-page":"177","DOI":"10.1016\/S0167-6393(99)00036-9","article-title":"Maximum likelihood modeling of pronunciation variation","volume":"29","author":"Holter","year":"1999","journal-title":"Speech Communication"},{"year":"1994","series-title":"Phonology in generative grammar","author":"Kenstowicz","key":"10.1016\/S0885-2308(03)00008-1_BIB11"},{"key":"10.1016\/S0885-2308(03)00008-1_BIB12","doi-asserted-by":"crossref","first-page":"193","DOI":"10.1016\/S0167-6393(99)00048-5","article-title":"Improving the performance of a Dutch CSR by modeling within-word and cross-word pronunciation variation","volume":"29","author":"Kessens","year":"1999","journal-title":"Speech Communication"},{"key":"10.1016\/S0885-2308(03)00008-1_BIB13","doi-asserted-by":"crossref","unstructured":"Li, A., Zheng, F., Byrne, W., Fung, P., Kamm, T.,Liu, Y., Song, Z., Ruhi, U., Venkataramani, V., Chen, X., 2000. CASS: a phonetically transcribed corpus of Mandarin spontaneous speech. In: Proceedings of the International Conference on Spoken Language Processing (ICSLP), Beijing, China","DOI":"10.21437\/ICSLP.2000-120"},{"key":"10.1016\/S0885-2308(03)00008-1_BIB14","doi-asserted-by":"crossref","unstructured":"Liu, Y., Fung, P., 1999. Decision tree-based triphones are robust and practical for Mandarin speech recognition. In: Proceedings of the European Conference on Speech Communication and Technology (Eurospeech), Budapest, Hungary, pp. 895\u2013898","DOI":"10.21437\/Eurospeech.1999-218"},{"key":"10.1016\/S0885-2308(03)00008-1_BIB15","doi-asserted-by":"crossref","unstructured":"Liu, Y., 2002. Pronunciation modeling for spontaneous Mandarin speech recognition. Ph.D. Thesis, The Hong Kong University of Science and Technology","DOI":"10.14711\/thesis-b775795"},{"key":"10.1016\/S0885-2308(03)00008-1_BIB16","doi-asserted-by":"crossref","unstructured":"Luo, X., Jelinek, F., 1999. Probabilistic classification of HMM states for large vocabulary continuous speech recognition. In: Proceedings of the IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP), Phoenix, USA, pp. 353\u2013356","DOI":"10.1109\/ICASSP.1999.758135"},{"year":"1993","series-title":"Fundamentals of speech recognition","author":"Rabiner","key":"10.1016\/S0885-2308(03)00008-1_BIB17"},{"key":"10.1016\/S0885-2308(03)00008-1_BIB18","doi-asserted-by":"crossref","first-page":"209","DOI":"10.1016\/S0167-6393(99)00037-0","article-title":"Stochastic pronunciation modelling from hand-labelled phonetic corpura","volume":"29","author":"Riley","year":"1999","journal-title":"Speech Communication"},{"key":"10.1016\/S0885-2308(03)00008-1_BIB19","series-title":"Automatic Speech and Speaker Recognition: Advanced Topics","first-page":"285","article-title":"Automatic generation of detailed pronunciation lexicions","author":"Riley","year":"1995"},{"key":"10.1016\/S0885-2308(03)00008-1_BIB20","doi-asserted-by":"crossref","first-page":"137","DOI":"10.1006\/csla.2000.0140","article-title":"Pronunciation modeling by sharing Gaussian densities across phonetic models","volume":"14","author":"Saraclar","year":"2000","journal-title":"Computer Speech and Language"},{"key":"10.1016\/S0885-2308(03)00008-1_BIB21","doi-asserted-by":"crossref","unstructured":"Saraclar, M., Khudanpur, S., 2000b. Pronunciation ambiguity vs pronunciation variability in speech recognition. In: Proceedings of the IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP), pp. 1679\u20131682","DOI":"10.1109\/ICASSP.2000.862073"},{"key":"10.1016\/S0885-2308(03)00008-1_BIB22","unstructured":"Saraclar, M., 2000c. Pronunciation modeling for conversational speech recognition. Ph.D. Thesis, The Johns Hopkins University, Baltimore, MD"},{"key":"10.1016\/S0885-2308(03)00008-1_BIB23","doi-asserted-by":"crossref","unstructured":"Sloboda, T., Waibel, A., 1996. Dictionary learning for spontaneous speech recognition. In: Proceedings of the International Conference on Spoken Language Processing (ICSLP), Philadelphia, USA, pp. 2328\u20132331","DOI":"10.1109\/ICSLP.1996.607274"},{"key":"10.1016\/S0885-2308(03)00008-1_BIB24","doi-asserted-by":"crossref","first-page":"225","DOI":"10.1016\/S0167-6393(99)00038-2","article-title":"Modeling pronunciation variation for ASR: a survey of the literature","volume":"29","author":"Strik","year":"1999","journal-title":"Speech Communication"},{"year":"1999","series-title":"The HTK Book","author":"Young","key":"10.1016\/S0885-2308(03)00008-1_BIB25"}],"container-title":["Computer Speech &amp; Language"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/api.elsevier.com\/content\/article\/PII:S0885230803000081?httpAccept=text\/xml","content-type":"text\/xml","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/api.elsevier.com\/content\/article\/PII:S0885230803000081?httpAccept=text\/plain","content-type":"text\/plain","content-version":"vor","intended-application":"text-mining"}],"deposited":{"date-parts":[[2024,12,11]],"date-time":"2024-12-11T23:34:14Z","timestamp":1733960054000},"score":1,"resource":{"primary":{"URL":"https:\/\/linkinghub.elsevier.com\/retrieve\/pii\/S0885230803000081"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2003,10]]},"references-count":25,"journal-issue":{"issue":"4","published-print":{"date-parts":[[2003,10]]}},"alternative-id":["S0885230803000081"],"URL":"https:\/\/doi.org\/10.1016\/s0885-2308(03)00008-1","relation":{},"ISSN":["0885-2308"],"issn-type":[{"type":"print","value":"0885-2308"}],"subject":[],"published":{"date-parts":[[2003,10]]}}}