{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,4]],"date-time":"2026-07-04T13:21:19Z","timestamp":1783171279857,"version":"3.54.6"},"reference-count":33,"publisher":"Springer Science and Business Media LLC","issue":"2","license":[{"start":{"date-parts":[[2026,3,28]],"date-time":"2026-03-28T00:00:00Z","timestamp":1774656000000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,3,28]],"date-time":"2026-03-28T00:00:00Z","timestamp":1774656000000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Evolving Systems"],"published-print":{"date-parts":[[2026,6]]},"DOI":"10.1007\/s12530-026-09817-x","type":"journal-article","created":{"date-parts":[[2026,3,28]],"date-time":"2026-03-28T05:23:12Z","timestamp":1774675392000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["Adapting language models with continual learning for temporal drifts"],"prefix":"10.1007","volume":"17","author":[{"given":"Antonio","family":"Carta","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Alberto Roberto","family":"Marinelli","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0003-4934-5344","authenticated-orcid":false,"given":"Lucia C.","family":"Passaro","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2026,3,28]]},"reference":[{"key":"9817_CR1","unstructured":"Bommasani R, Hudson DA, Adeli E, Altman RB, Arora S, Arx S, Bernstein MS, Bohg J, Bosselut A, Brunskill E, Brynjolfsson E, Buch S, Card D, Castellon R, Chatterji NS, Chen AS, Creel K, Davis JQ, Demszky D, Donahue C, Doumbouya M, Durmus E, Ermon S, Etchemendy J, Ethayarajh K, Fei-Fei L, Finn C, Gale T, Gillespie LE, Goel K, Goodman ND, Grossman S, Guha N, Hashimoto T, Henderson P, Hewitt J, Ho DE, Hong J, Hsu K, Huang J, Icard T, Jain S, Jurafsky D, Kalluri P, Karamcheti S, Keeling G, Khani F, Khattab O, Koh PW, Krass MS, Krishna R, Kuditipudi R, et al (2021) On the opportunities and risks of foundation models. https:\/\/arxiv.org\/abs\/2108.07258"},{"key":"9817_CR2","first-page":"9459","volume":"733","author":"P Lewis","year":"2020","unstructured":"Lewis P, Perez E, Piktus A, Petroni F, Karpukhin V, Goyal N, K\u00fcttler H, Lewis M, Yih W-T, Rockt\u00e4schel T (2020) Retrieval-augmented generation for knowledge-intensive nlp tasks. In: Proceedings of the 34th International Conference on Neural Information Processing Systems. Article No. 793, pp. 9459\u20139474","journal-title":"Proceedings of the 34th International Conference on Neural Information Processing Systems"},{"key":"9817_CR3","doi-asserted-by":"publisher","unstructured":"Jang J, Ye S, Lee C, Yang S, Shin J, Han J, Kim G, Seo M (2022) TemporalWiki: A Lifelong Benchmark for Training and Evaluating Ever-Evolving Language Models. In: Proceedings of the 2022 Conference on Empirical Methods in Natural Language Processing, Abu Dhabi, United Arab Emirates. Association for Computational Linguistics, pp. 6237\u20136250. https:\/\/doi.org\/10.18653\/v1\/2022.emnlp-main.418","DOI":"10.18653\/v1\/2022.emnlp-main.418"},{"issue":"4","key":"9817_CR4","doi-asserted-by":"publisher","first-page":"128","DOI":"10.1016\/S1364-6613(99)01294-2","volume":"3","author":"RM French","year":"1999","unstructured":"French RM (1999) Catastrophic forgetting in connectionist networks. Trends Cogn Sci 3(4):128\u2013135. https:\/\/doi.org\/10.1016\/S1364-6613(99)01294-2","journal-title":"Trends Cogn Sci"},{"key":"9817_CR5","unstructured":"Radford A, Wu J, Child R, Luan D, Amodei ea (2019) Language models are unsupervised multitask learners. OpenAI blog 1(8), 9"},{"key":"9817_CR6","doi-asserted-by":"publisher","unstructured":"Marinelli AR, Carta A, Passaro LC (2024) Updating knowledge in large language models: an empirical evaluation. In: 2024 IEEE International Conference on Evolving and Adaptive Intelligent Systems (EAIS), pp. 1\u20138. https:\/\/doi.org\/10.1109\/EAIS58494.2024.10570019","DOI":"10.1109\/EAIS58494.2024.10570019"},{"key":"9817_CR7","doi-asserted-by":"crossref","unstructured":"He T, Liu J, Cho K, Ott M, Liu B, et al (2021) Analyzing the forgetting problem in pretrain-finetuning of open-domain dialogue response models. In: Proceedings of the 16th Conference of the European Chapter of the Association for Computational Linguistics: Main Volume, pp. 1121\u20131133","DOI":"10.18653\/v1\/2021.eacl-main.95"},{"key":"9817_CR8","doi-asserted-by":"publisher","unstructured":"Li Y, Guerin F, Lin C (2024) Latesteval: addressing data contamination in language model evaluation through dynamic and time-sensitive test construction. In: Proceedings of the Thirty-Eighth AAAI Conference on Artificial Intelligence and Thirty-Sixth Conference on Innovative Applications of Artificial Intelligence and Fourteenth Symposium on Educational Advances in Artificial Intelligence. AAAI\u201924\/IAAI\u201924\/EAAI\u201924. AAAI Press (2024). https:\/\/doi.org\/10.1609\/aaai.v38i17.29822","DOI":"10.1609\/aaai.v38i17.29822"},{"key":"9817_CR9","doi-asserted-by":"crossref","unstructured":"Gururangan S, Marasovi\u0107 A, Swayamdipta S, Lo K, Beltagy I, Downey D, Smith NA (2020) Don\u2019t stop pretraining: adapt language models to domains and tasks. arXiv:2004.10964 [cs]","DOI":"10.18653\/v1\/2020.acl-main.740"},{"issue":"214","key":"9817_CR10","first-page":"1","volume":"24","author":"SV Mehta","year":"2023","unstructured":"Mehta SV, Patil D, Chandar S, Strubell E (2023) An empirical investigation of the role of pre-training in lifelong learning. J Mach Learn Res 24(214):1\u201350","journal-title":"J Mach Learn Res"},{"key":"9817_CR11","doi-asserted-by":"crossref","unstructured":"Jang J, Ye S, Lee C, Yang S, Shin J, et al (2023) TemporalWiki: A lifelong benchmark for training and evaluating ever-evolving language models","DOI":"10.18653\/v1\/2022.emnlp-main.418"},{"key":"9817_CR12","doi-asserted-by":"publisher","unstructured":"Cossu A, Carta A, Passaro L, Lomonaco V, Tuytelaars T, Bacciu D (2024) Continual pre-training mitigates forgetting in language and vision. Neural Networks 179:106492. https:\/\/doi.org\/10.1016\/j.neunet.2024.106492","DOI":"10.1016\/j.neunet.2024.106492"},{"key":"9817_CR13","unstructured":"Ke Z, Shao Y, Lin H, Konishi T, Kim G, Liu B (2023) Continual pre-training of language models. arXiv preprint arXiv:2302.03241"},{"key":"9817_CR14","doi-asserted-by":"crossref","unstructured":"Chen J, Chen Z, Wang J, Zhou K, Zhu Y, Jiang J, Min Y, Zhao X, Dou Z, Mao J, Lin Y, Song R, Xu J, Chen X, Yan R, Wei Z, Hu D, Huang W, Wen J-R (2025) Towards effective and efficient continual pre-training of large language models. In: Che W, Nabende J, Shutova E, Pilehvar MT (eds) Proceedings of the 63rd Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers), pp. 5779\u20135795. Association for Computational Linguistics, Vienna, Austria. https:\/\/aclanthology.org\/2025.acl-long.289\/","DOI":"10.18653\/v1\/2025.acl-long.289"},{"key":"9817_CR15","doi-asserted-by":"crossref","unstructured":"Zhang Z, Chen W, Lin Y, Wan H (2025) A generative adaptive replay continual learning model for temporal knowledge graph reasoning. In: Che W, Nabende J, Shutova E, Pilehvar MT (eds) Proceedings of the 63rd Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers), pp. 10964\u201310977. Association for Computational Linguistics, Vienna, Austria. https:\/\/aclanthology.org\/2025.acl-long.537\/","DOI":"10.18653\/v1\/2025.acl-long.537"},{"key":"9817_CR16","doi-asserted-by":"crossref","unstructured":"Li J, Armandpour M, Mirzadeh SI, Mehta S, Shankar V, Vemulapalli R, Bengio S, Tuzel O, Farajtabar M, Pouransari H, Faghri F (2025) TiC-LM: A web-scale benchmark for time-continual LLM pretraining. In: Che W, Nabende J, Shutova E, Pilehvar MT (eds) Proceedings of the 63rd Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers), pp. 32231\u201332273. Association for Computational Linguistics, Vienna, Austria. https:\/\/aclanthology.org\/2025.acl-long.1551\/","DOI":"10.18653\/v1\/2025.acl-long.1551"},{"key":"9817_CR17","doi-asserted-by":"publisher","unstructured":"Gupta K, Th\u00e9rien B, Ibrahim A, Richter ML, Anthony Q, Belilovsky E, Rish I, Lesort T (2023) Continual Pre-Training of Large Language Models: How to (Re)Warm Your Model? arXiv. https:\/\/doi.org\/10.48550\/arXiv.2308.04014","DOI":"10.48550\/arXiv.2308.04014"},{"key":"9817_CR18","doi-asserted-by":"publisher","unstructured":"Davari M, Asadi N, Mudur S, Aljundi R, Belilovsky E (2022) Probing representation forgetting in supervised and unsupervised continual learning. In: 2022 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR), pp. 16691\u201316700. https:\/\/doi.org\/10.1109\/CVPR52688.2022.01621","DOI":"10.1109\/CVPR52688.2022.01621"},{"key":"9817_CR19","unstructured":"Rolnick D, Ahuja A, Schwarz J, Lillicrap T, Wayne G (2019) Experience replay for continual learning. In: Wallach H, Larochelle H, Beygelzimer A, Alch\u00e9-Buc F, Fox E, Garnett R (eds) Advances in Neural Information Processing Systems, vol. 32. Curran Associates, Inc. https:\/\/proceedings.neurips.cc\/paper_files\/paper\/2019\/file\/fa7cdfad1a5aaf8370ebeda47a1ff1c3-Paper.pdf"},{"key":"9817_CR20","doi-asserted-by":"crossref","unstructured":"Huang J, Cui L, Wang A, Yang C, Liao X, Song L, Yao J, Su J (2024 ) Mitigating catastrophic forgetting in large language models with self-synthesized rehearsal. arXiv preprint https:\/\/arxiv.org\/pdf\/2403.01244","DOI":"10.18653\/v1\/2024.acl-long.77"},{"key":"9817_CR21","doi-asserted-by":"crossref","unstructured":"Li D, Chen Z, Cho E, Hao J, Liu X, Xing F, Guo C, Liu Y (2022) Overcoming catastrophic forgetting during domain adaptation of seq2seq language generation. In: Proceedings of the 2022 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies, pp. 5441\u20135454","DOI":"10.18653\/v1\/2022.naacl-main.398"},{"key":"9817_CR22","doi-asserted-by":"crossref","unstructured":"Qin Y, Zhang J, Lin Y, Liu Z, Li P, Sun M, Zhou J (2022) Elle: efficient lifelong pre-training for emerging data. arXiv preprint arXiv:2203.06311","DOI":"10.18653\/v1\/2022.findings-acl.220"},{"key":"9817_CR23","unstructured":"Masson\u00a0D\u2019Autume C, Ruder S, Kong L, Yogatama D (2019) Episodic memory in lifelong language learning. Adv Neural Inform Processing Syst 32"},{"key":"9817_CR24","first-page":"29348","volume":"34","author":"A Lazaridou","year":"2021","unstructured":"Lazaridou A, Kuncoro A, Gribovskaya E, Agrawal D, Liska A, Terzi T, Gimenez M, Masson d\u2019Autume C, Kocisky T, Ruder S (2021) Mind the gap: assessing temporal generalization in neural language models. Adv Neural Inf Process Syst 34:29348\u201329363","journal-title":"Adv Neural Inf Process Syst"},{"key":"9817_CR25","doi-asserted-by":"publisher","unstructured":"Madotto A, Lin Z, Zhou Z, Moon S, Crook P, Liu B, Yu Z, Cho E, Fung P, Wang Z (2021) Continual learning in task-oriented dialogue systems. In: Moens M-F, Huang X, Specia L, Yih SW-t (eds) Proceedings of the 2021 Conference on Empirical Methods in Natural Language Processing, pp. 7452\u20137467. Association for Computational Linguistics, Online and Punta Cana, Dominican Republic. https:\/\/doi.org\/10.18653\/v1\/2021.emnlp-main.590","DOI":"10.18653\/v1\/2021.emnlp-main.590"},{"key":"9817_CR26","doi-asserted-by":"crossref","unstructured":"Poth C, Sterz H, Paul I, Purkayastha S, Engl\u00e4nder L, Imhof T, Vuli\u0107 I, Ruder S, Gurevych I, Pfeiffer J (2023) Adapters: a unified library for parameter-efficient and modular transfer learning. https:\/\/arxiv.org\/abs\/2311.11077","DOI":"10.18653\/v1\/2023.emnlp-demo.13"},{"key":"9817_CR27","unstructured":"Parmar, J., Satheesh, S., Patwary, M., Shoeybi, M., Catanzaro, B.: Reuse, don\u2019t retrain: A recipe for continued pretraining of language models. arXiv preprint arXiv:2407.07263 (2024)"},{"key":"9817_CR28","unstructured":"Touvron H, Martin L, Stone K, Albert P, Almahairi A, Babaei Y, Bashlykov N, Batra S, Bhargava P, Bhosale S, et al. (2023) Llama 2: Open foundation and fine-tuned chat models. arXiv preprint arXiv:2307.09288"},{"key":"9817_CR29","unstructured":"Ye J, Chen X, Xu N, Zu C, Shao Z, Liu S, Cui Y, Zhou Z, Gong C, Shen Y, et al. (2023) A comprehensive capability analysis of gpt-3 and gpt-3.5 series models. arXiv preprint arXiv:2303.10420"},{"key":"9817_CR30","doi-asserted-by":"publisher","unstructured":"He T, Liu J, Cho K et al (2021) Analyzing the forgetting problem in pretrain-finetuning of open-domain dialogue response models. In: Proceedings of the 16th Conference of the European Chapter of the Association for Computational Linguistics: Main Volume, pp. 1121\u20131133. Association for Computational Linguistics, Online. https:\/\/doi.org\/10.18653\/v1\/2021.eacl-main.95","DOI":"10.18653\/v1\/2021.eacl-main.95"},{"key":"9817_CR31","unstructured":"Douze M, Guzhva A, Deng C, Johnson J, Szilvasy G, al. (2024) The faiss library arXiv:2401.08281 [cs.LG]"},{"key":"9817_CR32","doi-asserted-by":"crossref","unstructured":"Reimers N, Gurevych I (2019) Sentence-bert: Sentence embeddings using siamese bert-networks. In: Proceedings of the 2019 Conference on Empirical Methods in Natural Language Processing. Association for Computational Linguistics. https:\/\/arxiv.org\/abs\/1908.10084","DOI":"10.18653\/v1\/D19-1410"},{"key":"9817_CR33","unstructured":"Radford A, Wu J, Child R, Luan D, Amodei D, Sutskever I (2019) Language models are unsupervised multitask learners. https:\/\/api.semanticscholar.org\/CorpusID:160025533"}],"container-title":["Evolving Systems"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s12530-026-09817-x.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s12530-026-09817-x","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s12530-026-09817-x.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,7,4]],"date-time":"2026-07-04T12:57:38Z","timestamp":1783169858000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s12530-026-09817-x"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,3,28]]},"references-count":33,"journal-issue":{"issue":"2","published-print":{"date-parts":[[2026,6]]}},"alternative-id":["9817"],"URL":"https:\/\/doi.org\/10.1007\/s12530-026-09817-x","relation":{},"ISSN":["1868-6478","1868-6486"],"issn-type":[{"value":"1868-6478","type":"print"},{"value":"1868-6486","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,3,28]]},"assertion":[{"value":"27 February 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"5 August 2025","order":2,"name":"revised","label":"Revised","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"4 March 2026","order":3,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"28 March 2026","order":4,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors have no competing interests to declare that are relevant to the content of this article.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}],"article-number":"53"}}