{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,6,5]],"date-time":"2026-06-05T13:03:26Z","timestamp":1780664606635,"version":"3.54.1"},"reference-count":57,"publisher":"Springer Science and Business Media LLC","issue":"5","license":[{"start":{"date-parts":[[2026,3,30]],"date-time":"2026-03-30T00:00:00Z","timestamp":1774828800000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,3,30]],"date-time":"2026-03-30T00:00:00Z","timestamp":1774828800000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100003787","name":"Natural Science Foundation of Hebei Province","doi-asserted-by":"publisher","award":["F2024210008"],"award-info":[{"award-number":["F2024210008"]}],"id":[{"id":"10.13039\/501100003787","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Int. J. Mach. Learn. &amp; Cyber."],"published-print":{"date-parts":[[2026,5]]},"DOI":"10.1007\/s13042-025-02986-2","type":"journal-article","created":{"date-parts":[[2026,3,30]],"date-time":"2026-03-30T07:51:11Z","timestamp":1774857071000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["Pegasus-copynet: a novel summarization generation framework for scientific and technological texts"],"prefix":"10.1007","volume":"17","author":[{"given":"Shuhai","family":"Wang","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Haoran","family":"Wang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Xiangyang","family":"Wang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Yanmei","family":"Jiang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Shuo","family":"Sun","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Xiao","family":"Pan","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Peng","family":"Ren","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2026,3,30]]},"reference":[{"key":"2986_CR1","doi-asserted-by":"crossref","unstructured":"Sefid A, Giles CL (2022) SciBertSum: extractive summarization for scientific documents. In: International Workshop on Document Analysis Systems. Springer, Cham","DOI":"10.1007\/978-3-031-06555-2_46"},{"issue":"1","key":"2986_CR2","doi-asserted-by":"publisher","first-page":"164","DOI":"10.1177\/0165551521990616","volume":"49","author":"S Lamsiyah","year":"2023","unstructured":"Lamsiyah S, Mahdaouy AE, Ouatik SEA et al (2023) Unsupervised extractive multi-document summarization method based on transfer learning from bert multi-task fine-tuning. J Inf Sci 49(1):164\u2013182","journal-title":"J Inf Sci"},{"key":"2986_CR3","unstructured":"Zhang J, Zhao Y, Saleh M, et al. (2020) Pegasus: Pre-training with extracted gap-sentences for abstractive summarization. In: International Conference on Machine Learning. PMLR, 11328-11339"},{"key":"2986_CR4","doi-asserted-by":"crossref","unstructured":"Luo X, Li J, Chen Z (2023) Hgnn-t5 pegasus: A hybrid approach for chinese long text summarization. In: CCF Conference on Computer Supported Cooperative Work and Social Computing. Springer, 3-16","DOI":"10.1007\/978-981-99-9637-7_1"},{"key":"2986_CR5","doi-asserted-by":"crossref","unstructured":"Reddy RP, Charitha N, Gandla R, al (2023) An efficient abstractive summarization of research articles using pegasus model. In: International Conference on Information and Communication Technology for Intelligent Systems. Springer, 11-18","DOI":"10.1007\/978-981-99-3761-5_2"},{"key":"2986_CR6","doi-asserted-by":"crossref","unstructured":"Sarthak RV, Yadav P, al. (2023) Fine tuning the large language pegasus model for dialogue summarization. International Journal of Information Technology","DOI":"10.1007\/s41870-024-02307-w"},{"key":"2986_CR7","doi-asserted-by":"crossref","unstructured":"Vinitha M VS (2022) Review on recent advances in text summarization techniques. In: International Conference on Information and Management Engineering. Springer, 679-695","DOI":"10.1007\/978-981-99-2742-5_70"},{"key":"2986_CR8","doi-asserted-by":"crossref","unstructured":"Urlana\u00a0A RT, Mishra\u00a0P (2023) Controllable text summarization: Unraveling challenges, approaches, and prospects\u2013a survey. arXiv preprint arXiv:2311.09212","DOI":"10.18653\/v1\/2024.findings-acl.93"},{"key":"2986_CR9","doi-asserted-by":"crossref","unstructured":"Tewari\u00a0K KM, Yadav A K (2023) Extractive text summarization using statistical approach. In: Computer Vision and Machine Intelligence: Proceedings of CVMI 2022. Springer, 655-667","DOI":"10.1007\/978-981-19-7867-8_52"},{"key":"2986_CR10","doi-asserted-by":"crossref","unstructured":"Jaoua\u00a0M HAB (2003) Automatic text summarization of scientific articles based on classification of extract\u2019s population. In: International Conference on Intelligent Text Processing and Computational Linguistics, pp. 623\u2013634. Springer, Berlin, Heidelberg","DOI":"10.1007\/3-540-36456-0_70"},{"key":"2986_CR11","doi-asserted-by":"crossref","unstructured":"Beltagy I CA, Lo K (2019) Scibert: A pretrained language model for scientific text. arXiv preprint arXiv:1903.10676","DOI":"10.18653\/v1\/D19-1371"},{"issue":"1","key":"2986_CR12","first-page":"121","volume":"7","author":"NMMA Nazari","year":"2019","unstructured":"Nazari NMMA (2019) A survey on automatic text summarization. Journal of AI and Data Mining 7(1):121\u2013135","journal-title":"Journal of AI and Data Mining"},{"key":"2986_CR13","doi-asserted-by":"crossref","unstructured":"Ghorpade S CA, Khan\u00a0A al. (2022) A comparative analysis of textrank and lexrank algorithms using text summarization. In: International Joint Conference on Advances in Computational Intelligence, pp. 379\u2013393. Springer, Singapore","DOI":"10.1007\/978-981-97-0180-3_30"},{"key":"2986_CR14","doi-asserted-by":"crossref","unstructured":"Yang\u00a0Liu ML (2019) Text summarization with pretrained encoders. arXiv preprint arXiv:1908.08345","DOI":"10.18653\/v1\/D19-1387"},{"issue":"2","key":"2986_CR15","first-page":"306","volume":"41","author":"W Canyu","year":"2023","unstructured":"Canyu W, Xiaohai S, Yehui W, Rongbiao J, Yadong L, Shaoru Z, Shihao Y (2023) Generation method of extractive text summarization based on deep q-learning. Journal of Jilin University (Information Science Edition) 41(2):306\u2013314","journal-title":"Journal of Jilin University (Information Science Edition)"},{"key":"2986_CR16","doi-asserted-by":"crossref","unstructured":"Xu J CY, Gan Z al (2019) Discourse-aware neural extractive text summarization. arXiv preprint . arXiv:1910.14142","DOI":"10.18653\/v1\/2020.acl-main.451"},{"key":"2986_CR17","doi-asserted-by":"crossref","unstructured":"Yadav H JD, Patel N (2023) Fine-tuning bart for abstractive reviews summarization. In: Computational Intelligence: Select Proceedings of InCITe 2022, pp. 375\u2013385. Springer, Singapore","DOI":"10.1007\/978-981-19-7346-8_32"},{"key":"2986_CR18","doi-asserted-by":"crossref","unstructured":"See A MCD, Liu PJ (2017) Get to the point: Summarization with pointer-generator networks. arXiv preprint arXiv:1704.04368","DOI":"10.18653\/v1\/P17-1099"},{"key":"2986_CR19","unstructured":"Raffel C RA, Shazeer N, al. (2020) Exploring the limits of transfer learning with a unified text-to-text transformer. Journal of machine learning research 21(140), 1\u201367"},{"key":"2986_CR20","unstructured":"Dong L WW, Yang N, al. (2019) Unified language model pre-training for natural language understanding and generation. Advances in neural information processing systems 32"},{"key":"2986_CR21","doi-asserted-by":"crossref","unstructured":"Gururangan S SS, Marasovi\u0107 A al. (2020) Don\u2019t stop pretraining: Adapt language models to domains and tasks. arXiv preprint arXiv:2004.10964","DOI":"10.18653\/v1\/2020.acl-main.740"},{"key":"2986_CR22","doi-asserted-by":"crossref","unstructured":"Aghajanyan A SA, Gupta A, al. (2021) Muppet: Massive multi-task representations with pre-finetuning. arXiv preprint arXiv:2101.11038","DOI":"10.18653\/v1\/2021.emnlp-main.468"},{"issue":"7","key":"2986_CR23","doi-asserted-by":"publisher","first-page":"4115","DOI":"10.3390\/app13074115","volume":"13","author":"Y Gan","year":"2023","unstructured":"Gan Y, Lu G, Su Z, Wang L, Zhou J, Jiang J, Chen D (2023) A joint domain-specific pre-training method based on data enhancement. Appl Sci 13(7):4115","journal-title":"Appl Sci"},{"key":"2986_CR24","doi-asserted-by":"crossref","unstructured":"Huang\u00a0S XQ, Wang\u00a0R al. (2019) An extraction-abstraction hybrid approach for long document summarization. In: 2019 6th International Conference on Behavioral, Economic and Socio-Cultural Computing (BESC). IEEE, 1-6","DOI":"10.1109\/BESC48373.2019.8962979"},{"key":"2986_CR25","doi-asserted-by":"crossref","unstructured":"Hsu W\u00a0T LMY, Lin C\u00a0K, al. (2018) A unified model for extractive and abstractive summarization using inconsistency loss. arXiv preprint arXiv:1805.06266","DOI":"10.18653\/v1\/P18-1013"},{"key":"2986_CR26","unstructured":"Kaplan\u00a0J HT, McCandlish\u00a0S, al. (2020) Scaling laws for neural language models. arXiv preprint arXiv:2001.08361"},{"key":"2986_CR27","unstructured":"Li\u00a0W ZM, Peng\u00a0Y, al. (2023) Deep model fusion: A survey. arXiv preprint arXiv:2309.15698"},{"key":"2986_CR28","unstructured":"Sahoo\u00a0P SS, Singh A\u00a0K, al. (2024) A systematic survey of prompt engineering in large language models: Techniques and applications. arXiv preprint arXiv:2402.07927"},{"key":"2986_CR29","unstructured":"Chen\u00a0B LN, Zhang\u00a0Z, al. (2023) Unleashing the potential of prompt engineering in large language models: a comprehensive review. arXiv preprint arXiv:2310.14735"},{"key":"2986_CR30","doi-asserted-by":"crossref","unstructured":"Lin\u00a0X, Li Y, Wang W, al. (2024) Data-efficient fine-tuning for llm-based recommendation. In: Proceedings of the 47th International ACM SIGIR Conference on Research and Development in Information Retrieval, pp. 365\u2013374","DOI":"10.1145\/3626772.3657807"},{"key":"2986_CR31","doi-asserted-by":"crossref","unstructured":"Jacovi A, Shalom OS, Goldberg Y. (2018) Understanding convolutional neural networks for text classification. arXiv preprint arXiv:1809.08037","DOI":"10.18653\/v1\/W18-5408"},{"key":"2986_CR32","doi-asserted-by":"crossref","unstructured":"Takase S, Ri R, Kiyono S, al. (2024) Large vocabulary size improves large language models. arXiv preprint arXiv:2406.16508","DOI":"10.18653\/v1\/2025.findings-acl.57"},{"key":"2986_CR33","doi-asserted-by":"crossref","unstructured":"Xing\u00a0X, Zhang Q, Peng\u00a0M, al. (2020) Learning to generate representations for novel words: Mimic the oov situation in training. In: Natural Language Processing and Chinese Computing: 9th CCF International Conference, NLPCC 2020, Zhengzhou, China, October 14\u201318, 2020, Proceedings, Part I. Springer, 321-332","DOI":"10.1007\/978-3-030-60450-9_26"},{"key":"2986_CR34","doi-asserted-by":"crossref","unstructured":"Fetter\u00a0P, Kaltenmeier T, Kaltenmeier\u00a0A, al. (1996) Improved modeling of oov words in spontaneous speech. In: 1996 IEEE International Conference on Acoustics, Speech, and Signal Processing Conference Proceedings, vol. 1. IEEE, 534-537","DOI":"10.1109\/ICASSP.1996.541151"},{"key":"2986_CR35","doi-asserted-by":"crossref","unstructured":"Ren\u00a0F, Zhang D, Chen J (2024) An improved mt5 model for chinese text summary generation. In: CS IT Conference Proceedings, vol. 14","DOI":"10.5121\/csit.2024.140214"},{"key":"2986_CR36","unstructured":"Devlin J (2018) BERT: Pre-training of deep bidirectional transformers for language understanding. arXiv preprint arXiv:1810.04805"},{"key":"2986_CR37","doi-asserted-by":"crossref","unstructured":"Zeng\u00a0C, LS (2021) Analyzing the effect of masking length distribution of mlm: An evaluation framework and case study on chinese mrc datasets. Wireless Communications and Mobile Computing 2021(1), 5375334","DOI":"10.1155\/2021\/5375334"},{"key":"2986_CR38","doi-asserted-by":"publisher","first-page":"8518","DOI":"10.1109\/ACCESS.2022.3142178","volume":"10","author":"A Feng","year":"2022","unstructured":"Feng A, Zhang X, Song X (2022) Unrestricted attention may not be all you need-masked attention mechanism focuses better on relevant parts in aspect-based sentiment analysis. IEEE Access 10:8518\u20138528","journal-title":"IEEE Access"},{"key":"2986_CR39","doi-asserted-by":"crossref","unstructured":"Bedi P P\u00a0S, Sharma K. Bala\u00a0M (2023) Mlm: Masked language modeling using deep learning for efficient summarization of unstructured data. In: International Conference on Data Analytics Management, pp. 339\u2013347. Springer, Singapore","DOI":"10.1007\/978-981-99-6547-2_26"},{"key":"2986_CR40","doi-asserted-by":"crossref","unstructured":"Zou\u00a0D, Li X, Wang\u00a0S, al. (2024) Multispans: A multi-range spatial-temporal transformer network for traffic forecast via structural entropy optimization. In: Proceedings of the 17th ACM International Conference on Web Search and Data Mining, pp. 1032\u20131041","DOI":"10.1145\/3616855.3635820"},{"key":"2986_CR41","unstructured":"Cao\u00a0Y, L A. Peng\u00a0H, al. (2024) Multi-relational structural entropy. arXiv preprint 2405.07096"},{"key":"2986_CR42","doi-asserted-by":"crossref","unstructured":"Zou\u00a0D, H.X. Peng\u00a0H, al. (2023) Se-gsl: A general and effective graph structure learning framework through structural entropy optimization. In: Proceedings of the ACM Web Conference 2023, pp. 499\u2013510","DOI":"10.1145\/3543507.3583453"},{"issue":"5","key":"2986_CR43","doi-asserted-by":"publisher","first-page":"1618","DOI":"10.1016\/j.ipm.2019.05.003","volume":"56","author":"Z Kastrati","year":"2019","unstructured":"Kastrati Z, Yayilgan SY, Imran AS (2019) The impact of deep learning on document classification using semantically rich representations. Information Processing Management 56(5):1618\u20131632","journal-title":"Information Processing Management"},{"key":"2986_CR44","doi-asserted-by":"crossref","unstructured":"Plaza\u00a0L, C.-d.-A.J. (2013) Evaluating the use of different positional strategies for sentence selection in biomedical literature summarization. BMC Bioinformatics 14, 1\u201311","DOI":"10.1186\/1471-2105-14-71"},{"key":"2986_CR45","doi-asserted-by":"crossref","unstructured":"Gong\u00a0S, Qi J, Zhu\u00a0Z, al. (2022) Improving extractive document summarization with sentence centrality. PloS One 17(7), 0268278","DOI":"10.1371\/journal.pone.0268278"},{"key":"2986_CR46","doi-asserted-by":"crossref","unstructured":"Gu\u00a0J, Li H, Lu\u00a0Z, al. (2016) Incorporating copying mechanism in sequence-to-sequence learning. arXiv preprint arXiv:1603.06393","DOI":"10.18653\/v1\/P16-1154"},{"key":"2986_CR47","doi-asserted-by":"crossref","unstructured":"Zhang\u00a0Y, L.J. Wang\u00a0Y, al. (2018) A hierarchical attention seq2seq model with copynet for text summarization. In: 2018 International Conference on Robots Intelligent System (ICRIS). IEEE, 316-320","DOI":"10.1109\/ICRIS.2018.00086"},{"key":"2986_CR48","doi-asserted-by":"crossref","unstructured":"Peng\u00a0H, W.S. Li\u00a0J, al. (2019) Hierarchical taxonomy-aware and attentional graph capsule rcnns for large-scale multi-label text classification. IEEE Transactions on Knowledge and Data Engineering 33(6), 2505\u20132519","DOI":"10.1109\/TKDE.2019.2959991"},{"key":"2986_CR49","doi-asserted-by":"crossref","unstructured":"Peng\u00a0H, H.Y. Li\u00a0J, al. (32018) Large-scale hierarchical text classification with recursively regularized deep graph-cnn. In: Proceedings of the 2018 World Wide Web Conference, pp. 1063\u20131072","DOI":"10.1145\/3178876.3186005"},{"key":"2986_CR50","doi-asserted-by":"crossref","unstructured":"Guo\u00a0B, L.J. Zhang\u00a0C, al. (2019) Improving text classification with weighted word embeddings via a multi-channel textcnn model. Neurocomputing 363, 366\u2013374","DOI":"10.1016\/j.neucom.2019.07.052"},{"key":"2986_CR51","doi-asserted-by":"crossref","unstructured":"Li\u00a0H, M.X. Peng\u00a0Q, al. (2023) Abstractive financial news summarization via transformer-bilstm encoder and graph attention-based decoder. IEEE\/ACM Transactions on Audio, Speech, and Language Processing","DOI":"10.1109\/TASLP.2023.3304473"},{"key":"2986_CR52","doi-asserted-by":"crossref","unstructured":"Mulla\u00a0S, S.N.F. (2023) Weighted graph embedding feature with bi-directional long short-term memory classifier for multi-document text summarization. International Journal of Image and Graphics 24(02), 2450022","DOI":"10.1142\/S0219467824500220"},{"key":"2986_CR53","doi-asserted-by":"crossref","unstructured":"Banerjee\u00a0S, B.S. Mukherjee\u00a0S, al. (2023) An extract-then-abstract based method to generate disaster-news headlines using a dnn extractor followed by a transformer abstractor. Information Processing Management 60(3), 103291","DOI":"10.1016\/j.ipm.2023.103291"},{"key":"2986_CR54","doi-asserted-by":"crossref","unstructured":"Dilawari\u00a0A, S.S. Khan M U\u00a0G, al. (2023) Neural attention model for abstractive text summarization using linguistic feature space. IEEE Access 11, 23557\u201323564","DOI":"10.1109\/ACCESS.2023.3249783"},{"key":"2986_CR55","unstructured":"Yin\u00a0W, Y.M. Kann\u00a0K, al. (2017) Comparative study of cnn and rnn for natural language processing. arXiv preprint . arXiv:1702.01923"},{"key":"2986_CR56","unstructured":"Dong\u00a0L, W.W. Yang\u00a0N, al. (2019) Unified language model pre-training for natural language understanding and generation. Advances in Neural Information Processing Systems 32"},{"key":"2986_CR57","unstructured":"Raffel\u00a0C, R.A. Shazeer\u00a0N, al. (2020) Exploring the limits of transfer learning with a unified text-to-text transformer. Journal of Machine Learning Research 21(140), 1\u201367"}],"container-title":["International Journal of Machine Learning and Cybernetics"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s13042-025-02986-2.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s13042-025-02986-2","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s13042-025-02986-2.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,6,5]],"date-time":"2026-06-05T12:08:36Z","timestamp":1780661316000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s13042-025-02986-2"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,3,30]]},"references-count":57,"journal-issue":{"issue":"5","published-print":{"date-parts":[[2026,5]]}},"alternative-id":["2986"],"URL":"https:\/\/doi.org\/10.1007\/s13042-025-02986-2","relation":{},"ISSN":["1868-8071","1868-808X"],"issn-type":[{"value":"1868-8071","type":"print"},{"value":"1868-808X","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,3,30]]},"assertion":[{"value":"21 February 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"29 December 2025","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"30 March 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors have no relevant financial or non-financial interests to disclose. All authors certify that they have no affiliations with or involvement in any organization or entity with any financial interest or non-financial interest in the subject matter or materials discussed in this manuscript. The authors have no financial or proprietary interests in any material discussed in this article. The authors declare that they have no Conflict of interest.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}],"article-number":"230"}}