{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,6,2]],"date-time":"2026-06-02T13:04:47Z","timestamp":1780405487880,"version":"3.54.1"},"publisher-location":"Singapore","reference-count":34,"publisher":"Springer Nature Singapore","isbn-type":[{"value":"9789819688883","type":"print"},{"value":"9789819688890","type":"electronic"}],"license":[{"start":{"date-parts":[[2025,7,1]],"date-time":"2025-07-01T00:00:00Z","timestamp":1751328000000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2025,7,1]],"date-time":"2025-07-01T00:00:00Z","timestamp":1751328000000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":[],"published-print":{"date-parts":[[2026]]},"DOI":"10.1007\/978-981-96-8889-0_11","type":"book-chapter","created":{"date-parts":[[2025,6,30]],"date-time":"2025-06-30T08:56:10Z","timestamp":1751273770000},"page":"124-136","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":1,"title":["A Lightweight and Efficient Punctuation and Word Casing Prediction Model for On-Device Streaming ASR"],"prefix":"10.1007","author":[{"given":"Jian","family":"You","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Xiangfeng","family":"Li","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2025,7,1]]},"reference":[{"key":"11_CR1","doi-asserted-by":"crossref","unstructured":"He, Y., et al.: Streaming end-to-end speech recognition for mobile devices. In: ICASSP 2019\u20132019 IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP), pp. 6381\u20136385 (2019)","DOI":"10.1109\/ICASSP.2019.8682336"},{"key":"11_CR2","doi-asserted-by":"crossref","unstructured":"Mavandadi, S., Li, B., Zhang, C., Farris, B., Sainath, T.N. and Strohman, T.: A truly multilingual first pass and monolingual second pass streaming on-device ASR system. In: 2022 IEEE Spoken Language Technology Workshop (SLT), pp. 838\u2013845 (2022)","DOI":"10.1109\/SLT54892.2023.10023346"},{"key":"11_CR3","doi-asserted-by":"crossref","unstructured":"Jia, J., Li, K., Malek, M., Malik, K., Mahadeokar, J., Kalinli, O., Seide, F.: Joint federated learning and personalization for on-device ASR. In: 2023 IEEE Automatic Speech Recognition and Understanding Workshop (ASRU), pp. 1\u20138 (2023)","DOI":"10.1109\/ASRU57964.2023.10389738"},{"key":"11_CR4","unstructured":"Li, Y., et al.: Not All Weights Are Created Equal: Enhancing Energy Efficiency in On-Device Streaming Speech Recognition. arXiv preprint arXiv:2402.13076 (2024)"},{"key":"11_CR5","doi-asserted-by":"crossref","unstructured":"Szasz\u00e1k, G., T\u00fcndik, M.A.: Leveraging a character, word and prosody triplet for an ASR error robust and agglutination friendly punctuation approach. In: Interspeech, pp. 2988\u20132992 (2019)","DOI":"10.21437\/Interspeech.2019-2132"},{"key":"11_CR6","doi-asserted-by":"crossref","unstructured":"Che, X., Luo, S., Yang, H., Meinel, C.: Sentence boundary detection based on parallel lexical and acoustic models. In:\u00a0Interspeech, pp. 2528\u20132532 (2016)","DOI":"10.21437\/Interspeech.2016-257"},{"key":"11_CR7","doi-asserted-by":"crossref","unstructured":"Xu, C., Xie, L., Huang, G., Xiao, X., Chng, E., Li, H.: A deep neural network approach for sentence boundary detection in broadcast news. In: Interspeech, pp. 2887\u20132891 (2014)","DOI":"10.21437\/Interspeech.2014-599"},{"key":"11_CR8","doi-asserted-by":"crossref","unstructured":"Chen, Q., Chen, M., Li, B., Wang, W.: Controllable time-delay transformer for real-time punctuation prediction and disfluency detection. In: ICASSP 2020\u20132020 IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP), pp. 8069\u20138073 (2020)","DOI":"10.1109\/ICASSP40776.2020.9053159"},{"key":"11_CR9","doi-asserted-by":"publisher","unstructured":"Tran, H., Dinh, C.V., Pham, Q. and Nguyen, B.T.: An efficient transformer-based model for Vietnamese punctuation prediction. In: Fujita, H., Selamat, A., Lin, J.CW., Ali, M. (eds.) IEA\/AIE 2021. LNCS, vol. 12799, pp. 47\u201358. Springer, Cham (2021). https:\/\/doi.org\/10.1007\/978-3-030-79463-7_5","DOI":"10.1007\/978-3-030-79463-7_5"},{"key":"11_CR10","doi-asserted-by":"crossref","unstructured":"Huang, Q., Ko, T., Tang, H.L., Liu, X., Wu, B.: Token-level supervised contrastive learning for punctuation restoration. arXiv preprint arXiv:2107.09099 (2021)","DOI":"10.21437\/Interspeech.2021-661"},{"key":"11_CR11","unstructured":"Nagy, A., Bial, B. and \u00c1cs, J.: Automatic punctuation restoration with bert models. arXiv preprint arXiv:2101.07343 (2021)"},{"key":"11_CR12","unstructured":"Lu, W., Ng, H.T.: Better punctuation prediction with dynamic conditional random fields. In: Proceedings of the 2010 Conference on Empirical Methods in Natural Language Processing, pp. 177\u2013186 (2010)"},{"key":"11_CR13","unstructured":"Zhang, D., Wu, S., Yang, N., Li, M.: Punctuation prediction with transition-based parsing. In: Proceedings of the 51st Annual Meeting of the Association for Computational Linguistics, pp. 752\u2013760 (2013)"},{"key":"11_CR14","unstructured":"Che, X., Wang, C., Yang, H., Meinel, C.: Punctuation prediction for unsegmented transcript based on word vector. In: Proceedings of the Tenth International Conference on Language Resources and Evaluation (LREC\u201916), pp. 654\u2013658 (2016)"},{"key":"11_CR15","doi-asserted-by":"crossref","unstructured":"\u017belasko, P., Szyma\u0144ski, P., Mizgajski, J., Szymczak, A., Carmiel, Y., Dehak, N.: Punctuation prediction model for conversational speech. arXiv preprint arXiv:1807.00543 (2018)","DOI":"10.21437\/Interspeech.2018-1096"},{"key":"11_CR16","doi-asserted-by":"crossref","unstructured":"Augustyniak, \u0141., et al.: Punctuation prediction in spontaneous conversations: Can we mitigate asr errors with retrofitted word embeddings? In: Interspeech, pp. 4906\u20134910 (2020)","DOI":"10.21437\/Interspeech.2020-1250"},{"key":"11_CR17","doi-asserted-by":"crossref","unstructured":"Tilk, O., Alum\u00e4e, T.: Bidirectional recurrent neural network with attention mechanism for punctuation restoration. In: Interspeech, pp. 3047\u20133051 (2016)","DOI":"10.21437\/Interspeech.2016-1517"},{"key":"11_CR18","doi-asserted-by":"crossref","unstructured":"T\u00fcndik, M.\u00c1., Tarj\u00e1n, B., Szasz\u00e1k, G.: Low Latency MaxEnt-and RNN-based word sequence models for punctuation restoration of closed caption data. In: Statistical Language and Speech Processing: 5th International Conference, SLSP 2017, pp. 155\u2013166 (2017)","DOI":"10.1007\/978-3-319-68456-7_13"},{"key":"11_CR19","doi-asserted-by":"crossref","unstructured":"\u00d6ktem, A., Farr\u00fas, M., Wanner, L.: Attentional parallel RNNs for generating punctuation in transcribed speech. In: Statistical Language and Speech Processing: 5th International Conference, SLSP 2017, pp. 131\u2013142 (2017)","DOI":"10.1007\/978-3-319-68456-7_11"},{"key":"11_CR20","doi-asserted-by":"crossref","unstructured":"Pahuja, V., Laha, A., Mirkin, S., Raykar, V., Kotlerman, L., Lev, G.: Joint learning of correlated sequence labelling tasks using bidirectional recurrent neural networks. In: Interspeech, pp. 548\u2013552 (2017)","DOI":"10.21437\/Interspeech.2017-1247"},{"key":"11_CR21","doi-asserted-by":"publisher","DOI":"10.1016\/j.petrol.2021.108838","volume":"205","author":"L Shan","year":"2021","unstructured":"Shan, L., Liu, Y., Tang, M., Yang, M., Bai, X.: CNN-BiLSTM hybrid neural networks with attention mechanism for well log prediction. J. Petrol. Sci. Eng. 205, 108838 (2021)","journal-title":"J. Petrol. Sci. Eng."},{"key":"11_CR22","doi-asserted-by":"publisher","first-page":"86","DOI":"10.1007\/978-3-030-83527-9_7","volume-title":"Text, Speech, and Dialogue: 24th International Conference, TSD 2021, Olomouc, Czech Republic, 6\u20139 Sep 2021, Proceedings","author":"J \u0160vec","year":"2021","unstructured":"\u0160vec, J., Lehe\u010dka, J., \u0160m\u00eddl, L., Ircing, P.: Transformer-based automatic punctuation prediction and word casing reconstruction of the ASR output. In: Ek\u0161tein, K., P\u00e1rtl, F., Konop\u00edk, M. (eds.) Text, Speech, and Dialogue: 24th International Conference, TSD 2021, Olomouc, Czech Republic, 6\u20139 Sep 2021, Proceedings, pp. 86\u201394. Springer International Publishing, Cham (2021). https:\/\/doi.org\/10.1007\/978-3-030-83527-9_7"},{"key":"11_CR23","doi-asserted-by":"crossref","unstructured":"Dixit, M.S.S.R.K., Kirchhoff, S.B.K.: Robust Prediction of Punctuation and Truecasing for Medical ASR. In: Proceedings of the First Workshop on Natural Language Processing for Medical Conversations, pp. 53\u201362 (2020)","DOI":"10.18653\/v1\/2020.nlpmc-1.8"},{"key":"11_CR24","doi-asserted-by":"crossref","unstructured":"Alam, T., Khan, A.. Alam, F.: Punctuation restoration using transformer models for high-and low-resource languages. In: Proceedings of the Sixth Workshop on Noisy User-generated Text (W-NUT 2020), pp. 132\u2013142 (2020)","DOI":"10.18653\/v1\/2020.wnut-1.18"},{"key":"11_CR25","unstructured":"Federico, M., Bentivogli, L., Paul, M., St\u00fcker, S.: Overview of the IWSLT 2011 evaluation campaign. In: 2011 international workshop on spoken language translation, IWSLT 2011, pp. 11\u201327 (2011)"},{"key":"11_CR26","doi-asserted-by":"crossref","unstructured":"Kudo, T., Richardson, J.: Sentencepiece: A simple and language independent subword tokenizer and detokenizer for neural text processing. arXiv preprint arXiv:1808.06226 (2018)","DOI":"10.18653\/v1\/D18-2012"},{"key":"11_CR27","unstructured":"Paszke, A., et al.: Pytorch: An imperative style, high-performance deep learning library. In: Advances in neural information processing systems (NeurIPS), pp. 8024\u20138035 (2019)"},{"key":"11_CR28","doi-asserted-by":"crossref","unstructured":"He, K., Zhang, X., Ren, S., Sun, J.: Delving deep into rectifiers: surpassing human-level performance on imagenet classification. In: Proceedings of the IEEE International Conference on Computer Vision, pp. 1026\u20131034 (2015)","DOI":"10.1109\/ICCV.2015.123"},{"key":"11_CR29","unstructured":"Kingma, D.P.: Adam: A method for stochastic optimization. arXiv preprint arXiv:1412.6980 (2014)"},{"key":"11_CR30","doi-asserted-by":"crossref","unstructured":"Tilk, O., Alum\u00e4e, T.: LSTM for punctuation restoration in speech transcripts. In: Interspeech, pp. 683\u2013687 (2015)","DOI":"10.21437\/Interspeech.2015-240"},{"key":"11_CR31","doi-asserted-by":"crossref","unstructured":"Yi, J., Tao, J., Wen, Z., Li, Y.: Distilling knowledge from an ensemble of models for punctuation prediction. In: Interspeech, pp. 2779\u20132783 (2017)","DOI":"10.21437\/Interspeech.2017-1079"},{"key":"11_CR32","doi-asserted-by":"crossref","unstructured":"Yi, J. and Tao, J.: Self-attention based model for punctuation prediction using word and speech embeddings. In: ICASSP 2019\u20132019 IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP), pp. 7270\u20137274 (2019)","DOI":"10.1109\/ICASSP.2019.8682260"},{"key":"11_CR33","unstructured":"Vaswani, A., et al.: Attention is all you need. In: Advances in neural information processing systems (NeurIPS), pp. 5998\u20136008 (2017)"},{"key":"11_CR34","unstructured":"Open neural network exchange (onnx). https:\/\/github.com\/microsoft\/onnxruntime. Last accessed 2024"}],"container-title":["Lecture Notes in Computer Science","Advances and Trends in Artificial Intelligence. Theory and Applications"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/978-981-96-8889-0_11","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,4,1]],"date-time":"2026-04-01T01:34:35Z","timestamp":1775007275000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/978-981-96-8889-0_11"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,7,1]]},"ISBN":["9789819688883","9789819688890"],"references-count":34,"URL":"https:\/\/doi.org\/10.1007\/978-981-96-8889-0_11","relation":{},"ISSN":["0302-9743","1611-3349"],"issn-type":[{"value":"0302-9743","type":"print"},{"value":"1611-3349","type":"electronic"}],"subject":[],"published":{"date-parts":[[2025,7,1]]},"assertion":[{"value":"1 July 2025","order":1,"name":"first_online","label":"First Online","group":{"name":"ChapterHistory","label":"Chapter History"}},{"value":"The authors have no competing interests to declare that are relevant to the content of this article.","order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Disclosure of Interests"}},{"value":"IEA\/AIE","order":1,"name":"conference_acronym","label":"Conference Acronym","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"International Conference on Industrial, Engineering and Other Applications of Applied Intelligent Systems","order":2,"name":"conference_name","label":"Conference Name","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"Kytakyushu","order":3,"name":"conference_city","label":"Conference City","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"Japan","order":4,"name":"conference_country","label":"Conference Country","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"2025","order":5,"name":"conference_year","label":"Conference Year","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"1 July 2025","order":7,"name":"conference_start_date","label":"Conference Start Date","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"4 July 2025","order":8,"name":"conference_end_date","label":"Conference End Date","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"38","order":9,"name":"conference_number","label":"Conference Number","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"ieaaie2025","order":10,"name":"conference_id","label":"Conference ID","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"https:\/\/www.i-somet.org\/iea-aie2025\/committees.html","order":11,"name":"conference_url","label":"Conference URL","group":{"name":"ConferenceInfo","label":"Conference Information"}}]}}