{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,4,1]],"date-time":"2026-04-01T16:03:54Z","timestamp":1775059434921,"version":"3.50.1"},"publisher-location":"Cham","reference-count":33,"publisher":"Springer Nature Switzerland","isbn-type":[{"value":"9783031705458","type":"print"},{"value":"9783031705465","type":"electronic"}],"license":[{"start":{"date-parts":[[2024,1,1]],"date-time":"2024-01-01T00:00:00Z","timestamp":1704067200000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2024,1,1]],"date-time":"2024-01-01T00:00:00Z","timestamp":1704067200000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":[],"published-print":{"date-parts":[[2024]]},"DOI":"10.1007\/978-3-031-70546-5_15","type":"book-chapter","created":{"date-parts":[[2024,9,10]],"date-time":"2024-09-10T05:02:47Z","timestamp":1725944567000},"page":"253-269","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":2,"title":["Global-SEG: Text Semantic Segmentation Based on\u00a0Global Semantic Pair Relations"],"prefix":"10.1007","author":[{"ORCID":"https:\/\/orcid.org\/0009-0002-7857-8737","authenticated-orcid":false,"given":"Wenjun","family":"Sun","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-5993-1630","authenticated-orcid":false,"given":"Hanh Thi Hong","family":"Tran","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-0787-2990","authenticated-orcid":false,"given":"Carlos-Emiliano","family":"Gonz\u00e1lez-Gallardo","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-0123-439X","authenticated-orcid":false,"given":"Micka\u00ebl","family":"Coustaty","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0001-6160-3356","authenticated-orcid":false,"given":"Antoine","family":"Doucet","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2024,9,11]]},"reference":[{"key":"15_CR1","doi-asserted-by":"publisher","first-page":"169","DOI":"10.1162\/tacl_a_00261","volume":"7","author":"S Arnold","year":"2019","unstructured":"Arnold, S., Schneider, R., Cudr\u00e9-Mauroux, P., Gers, F.A., L\u00f6ser, A.: Sector: a neural model for coherent topic segmentation and classification. Trans. Assoc. Comput. Linguist. 7, 169\u2013184 (2019)","journal-title":"Trans. Assoc. Comput. Linguist."},{"key":"15_CR2","doi-asserted-by":"crossref","unstructured":"Barrow, J., Jain, R., Morariu, V., Manjunatha, V., Oard, D.W., Resnik, P.: A joint model for document segmentation and segment labeling. In: Proceedings of the 58th Annual Meeting of the Association for Computational Linguistics, pp. 313\u2013322 (2020)","DOI":"10.18653\/v1\/2020.acl-main.29"},{"key":"15_CR3","doi-asserted-by":"publisher","first-page":"177","DOI":"10.1023\/A:1007506220214","volume":"34","author":"D Beeferman","year":"1999","unstructured":"Beeferman, D., Berger, A., Lafferty, J.: Statistical models for text segmentation. Mach. Learn. 34, 177\u2013210 (1999)","journal-title":"Mach. Learn."},{"key":"15_CR4","doi-asserted-by":"publisher","first-page":"135","DOI":"10.1162\/tacl_a_00051","volume":"5","author":"P Bojanowski","year":"2017","unstructured":"Bojanowski, P., Grave, E., Joulin, A., Mikolov, T.: Enriching word vectors with subword information. Trans. Assoc. Comput. Linguist. 5, 135\u2013146 (2017)","journal-title":"Trans. Assoc. Comput. Linguist."},{"key":"15_CR5","doi-asserted-by":"crossref","unstructured":"Boro\u015f, E., et al.: Alleviating digitization errors in named entity recognition for historical documents. In: Proceedings of the 24th Conference on Computational Natural Language Learning, pp. 431\u2013441 (2020)","DOI":"10.18653\/v1\/2020.conll-1.35"},{"key":"15_CR6","doi-asserted-by":"crossref","unstructured":"Chen, H., Branavan, S., Barzilay, R., Karger, D.R.: Global models of document structure using latent permutations. In: Proceedings of Human Language Technologies: The 2009 Annual Conference of the North American Chapter of the Association for Computational Linguistics, pp. 371\u2013379. Association for Computational Linguistics (2009)","DOI":"10.3115\/1620754.1620808"},{"key":"15_CR7","doi-asserted-by":"crossref","unstructured":"Conneau, A., Kiela, D., Schwenk, H., Barrault, L., Bordes, A.: Supervised learning of universal sentence representations from natural language inference data. In: Proceedings of the 2017 Conference on Empirical Methods in Natural Language Processing, pp. 670\u2013680. Association for Computational Linguistics (2017)","DOI":"10.18653\/v1\/D17-1070"},{"key":"15_CR8","doi-asserted-by":"publisher","unstructured":"Devlin, J., Chang, M.W., Lee, K., Toutanova, K.: BERT: pre-training of deep bidirectional transformers for language understanding. In: Proceedings of the 2019 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies, Minneapolis, Minnesota, Volume 1 (Long and Short Papers), pp. 4171\u20134186. Association for Computational Linguistics, June 2019. https:\/\/doi.org\/10.18653\/v1\/N19-1423. https:\/\/aclanthology.org\/N19-1423","DOI":"10.18653\/v1\/N19-1423"},{"key":"15_CR9","doi-asserted-by":"crossref","unstructured":"Ehrmann, M., et al.: Extended overview of HIPE-2022: named entity recognition and linking in multilingual historical documents. In: CEUR Workshop Proceedings, pp. 1038\u20131063. No.\u00a03180. CEUR-WS (2022)","DOI":"10.1007\/978-3-031-13643-6_26"},{"key":"15_CR10","doi-asserted-by":"crossref","unstructured":"Gao, T., Yao, X., Chen, D.: SimCSE: simple contrastive learning of sentence embeddings. In: 2021 Conference on Empirical Methods in Natural Language Processing, EMNLP 2021, pp. 6894\u20136910. Association for Computational Linguistics (ACL) (2021)","DOI":"10.18653\/v1\/2021.emnlp-main.552"},{"key":"15_CR11","series-title":"Lecture Notes in Computer Science","doi-asserted-by":"publisher","first-page":"76","DOI":"10.1007\/978-981-99-8085-7_7","volume-title":"ICADL 2023","author":"N Girdhar","year":"2023","unstructured":"Girdhar, N., Coustaty, M., Doucet, A.: Benchmarking NAS for article separation in historical newspapers. In: Goh, D.H., Chen, S.J., Tuarob, S. (eds.) ICADL 2023. LNCS, vol. 14457, pp. 76\u201388. Springer, Singapore (2023). https:\/\/doi.org\/10.1007\/978-981-99-8085-7_7"},{"key":"15_CR12","doi-asserted-by":"crossref","unstructured":"Glava\u0161, G., Nanni, F., Ponzetto, S.P.: Unsupervised text segmentation using semantic relatedness graphs. In: Proceedings of the Fifth Joint Conference on Lexical and Computational Semantics, pp. 125\u2013130. Association for Computational Linguistics (2016)","DOI":"10.18653\/v1\/S16-2016"},{"key":"15_CR13","doi-asserted-by":"crossref","unstructured":"Glava\u0161, G., Somasundaran, S.: Two-level transformer and auxiliary coherence modeling for improved text segmentation. In: Proceedings of the AAAI Conference on Artificial Intelligence, vol. 34, no. 05, pp. 7797\u20137804 (2020)","DOI":"10.1609\/aaai.v34i05.6284"},{"key":"15_CR14","series-title":"Lecture Notes in Computer Science","doi-asserted-by":"publisher","first-page":"213","DOI":"10.1007\/978-3-031-00129-1_14","volume-title":"DASFAA 2022, Part III","author":"Z Gong","year":"2022","unstructured":"Gong, Z., et al.: Tipster: a topic-guided language model for topic-aware text segmentation. In: Bhattacharya, A., et al. (eds.) DASFAA 2022, Part III. Lecture Notes in Computer Science, vol. 13247, pp. 213\u2013221. Springer, Cham (2022). https:\/\/doi.org\/10.1007\/978-3-031-00129-1_14"},{"key":"15_CR15","doi-asserted-by":"crossref","unstructured":"Hearst, M.A.: Multi-paragraph segmentation expository text. In: 32nd Annual Meeting of the Association for Computational Linguistics, pp. 9\u201316 (1994)","DOI":"10.3115\/981732.981734"},{"key":"15_CR16","doi-asserted-by":"crossref","unstructured":"Li, B., Zhou, H., He, J., Wang, M., Yang, Y., Li, L.: On the sentence embeddings from pre-trained language models. In: Proceedings of the 2020 Conference on Empirical Methods in Natural Language Processing (EMNLP), pp. 9119\u20139130 (2020)","DOI":"10.18653\/v1\/2020.emnlp-main.733"},{"key":"15_CR17","doi-asserted-by":"crossref","unstructured":"Li, J., Sun, A., Joty, S.R.: SegBot: a generic neural text segmentation model with pointer network. In: IJCAI, pp. 4166\u20134172 (2018)","DOI":"10.24963\/ijcai.2018\/579"},{"key":"15_CR18","doi-asserted-by":"crossref","unstructured":"Lo, K., Jin, Y., Tan, W., Liu, M., Du, L., Buntine, W.: Transformer over pre-trained transformer for neural text segmentation with enhanced topic coherence. In: Findings of the Association for Computational Linguistics: EMNLP 2021, pp. 3334\u20133340 (2021)","DOI":"10.18653\/v1\/2021.findings-emnlp.283"},{"key":"15_CR19","doi-asserted-by":"publisher","unstructured":"Lukasik, M., Dadachev, B., Papineni, K., Sim\u00f5es, G.: Text segmentation by cross segment attention. In: Proceedings of the 2020 Conference on Empirical Methods in Natural Language Processing (EMNLP), pp. 4707\u20134716. Association for Computational Linguistics, Online, November 2020. https:\/\/doi.org\/10.18653\/v1\/2020.emnlp-main.380. https:\/\/aclanthology.org\/2020.emnlp-main.380","DOI":"10.18653\/v1\/2020.emnlp-main.380"},{"key":"15_CR20","doi-asserted-by":"crossref","unstructured":"Moro, G., Ragazzi, L.: Semantic self-segmentation for abstractive summarization of long documents in low-resource regimes. In: Proceedings of the AAAI Conference on Artificial Intelligence, vol.\u00a036, pp. 11085\u201311093 (2022)","DOI":"10.1609\/aaai.v36i10.21357"},{"key":"15_CR21","unstructured":"OpenAI: GPT-4 technical report (2023)"},{"key":"15_CR22","unstructured":"Radford, A., Narasimhan, K., Salimans, T., Sutskever, I., et\u00a0al.: Improving language understanding by generative pre-training (2018)"},{"key":"15_CR23","doi-asserted-by":"crossref","unstructured":"Reimers, N., Gurevych, I.: Sentence-BERT: sentence embeddings using Siamese BERT-networks. In: Proceedings of the 2019 Conference on Empirical Methods in Natural Language Processing and the 9th International Joint Conference on Natural Language Processing (EMNLP-IJCNLP), pp. 3982\u20133992 (2019)","DOI":"10.18653\/v1\/D19-1410"},{"key":"15_CR24","unstructured":"Riedl, M., Biemann, C.: TopicTiling: a text segmentation algorithm based on LDA. In: Proceedings of ACL 2012 Student Research Workshop, pp. 37\u201342 (2012)"},{"key":"15_CR25","unstructured":"Schweter, S., M\u00e4rz, L., Schmid, K., \u00c7ano, E.: hmBERT: historical multilingual language models for named entity recognition. In: Proceedings of the Working Notes of CLEF 2022 - Conference and Labs of the Evaluation Forum, vol. 3180, pp. 1109\u20131129, September 2022. http:\/\/eprints.cs.univie.ac.at\/7549\/"},{"key":"15_CR26","unstructured":"Touvron, H., et\u00a0al.: LLaMA: open and efficient foundation language models. arXiv preprint arXiv:2302.13971 (2023)"},{"key":"15_CR27","unstructured":"Touvron, H., et\u00a0al.: LLaMA 2: open foundation and fine-tuned chat models. arXiv preprint arXiv:2307.09288 (2023)"},{"key":"15_CR28","doi-asserted-by":"crossref","unstructured":"Utiyama, M., Isahara, H.: A statistical model for domain-independent text segmentation. In: Proceedings of the 39th Annual Meeting of the Association for Computational Linguistics, pp. 499\u2013506 (2001)","DOI":"10.3115\/1073012.1073076"},{"key":"15_CR29","unstructured":"Vaswani, A., et al.: Attention is all you need. In: Advances in Neural Information Processing Systems, vol.\u00a030, pp. 5998\u20136008. https:\/\/proceedings.neurips.cc\/paper_files\/paper\/2017\/file\/3f5ee243547dee91fbd053c1c4a845aa-Paper.pdf"},{"key":"15_CR30","doi-asserted-by":"crossref","unstructured":"Wang, L., Li, S., L\u00fc, Y., Wang, H.: Learning to rank semantic coherence for topic segmentation. In: Proceedings of the 2017 Conference on Empirical Methods in Natural Language Processing, pp. 1340\u20131344 (2017)","DOI":"10.18653\/v1\/D17-1139"},{"key":"15_CR31","doi-asserted-by":"crossref","unstructured":"Xia, J., et al.: Dialogue topic segmentation via parallel extraction network with neighbor smoothing. In: Proceedings of the 45th International ACM SIGIR Conference on Research and Development in Information Retrieval, pp. 2126\u20132131 (2022)","DOI":"10.1145\/3477495.3531817"},{"key":"15_CR32","doi-asserted-by":"crossref","unstructured":"Xu, Y., Li, M., Cui, L., Huang, S., Wei, F., Zhou, M.: LayoutLM: pre-training of text and layout for document image understanding. In: Proceedings of the 26th ACM SIGKDD International Conference on Knowledge Discovery & Data Mining, pp. 1192\u20131200 (2020)","DOI":"10.1145\/3394486.3403172"},{"key":"15_CR33","doi-asserted-by":"crossref","unstructured":"Zhang, N., et al.: Document-level relation extraction as semantic segmentation. In: Proceedings of the Thirtieth International Joint Conference on Artificial Intelligence, pp. 3999\u20134006 (2021)","DOI":"10.24963\/ijcai.2021\/551"}],"container-title":["Lecture Notes in Computer Science","Document Analysis and Recognition - ICDAR 2024"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/978-3-031-70546-5_15","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2024,9,10]],"date-time":"2024-09-10T05:06:52Z","timestamp":1725944812000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/978-3-031-70546-5_15"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2024]]},"ISBN":["9783031705458","9783031705465"],"references-count":33,"URL":"https:\/\/doi.org\/10.1007\/978-3-031-70546-5_15","relation":{},"ISSN":["0302-9743","1611-3349"],"issn-type":[{"value":"0302-9743","type":"print"},{"value":"1611-3349","type":"electronic"}],"subject":[],"published":{"date-parts":[[2024]]},"assertion":[{"value":"11 September 2024","order":1,"name":"first_online","label":"First Online","group":{"name":"ChapterHistory","label":"Chapter History"}},{"value":"ICDAR","order":1,"name":"conference_acronym","label":"Conference Acronym","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"International Conference on Document Analysis and Recognition","order":2,"name":"conference_name","label":"Conference Name","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"Athens","order":3,"name":"conference_city","label":"Conference City","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"Greece","order":4,"name":"conference_country","label":"Conference Country","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"2024","order":5,"name":"conference_year","label":"Conference Year","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"30 August 2024","order":7,"name":"conference_start_date","label":"Conference Start Date","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"4 September 2024","order":8,"name":"conference_end_date","label":"Conference End Date","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"icdar2024","order":10,"name":"conference_id","label":"Conference ID","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"https:\/\/icdar2024.net\/","order":11,"name":"conference_url","label":"Conference URL","group":{"name":"ConferenceInfo","label":"Conference Information"}}]}}