{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,3,31]],"date-time":"2026-03-31T06:42:36Z","timestamp":1774939356297,"version":"3.50.1"},"reference-count":38,"publisher":"Springer Science and Business Media LLC","issue":"4","license":[{"start":{"date-parts":[[2026,2,1]],"date-time":"2026-02-01T00:00:00Z","timestamp":1769904000000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,2,1]],"date-time":"2026-02-01T00:00:00Z","timestamp":1769904000000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"publisher","award":["61906021"],"award-info":[{"award-number":["61906021"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Appl Intell"],"published-print":{"date-parts":[[2026,2]]},"DOI":"10.1007\/s10489-026-07090-8","type":"journal-article","created":{"date-parts":[[2026,2,26]],"date-time":"2026-02-26T13:19:25Z","timestamp":1772111965000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["MQSF-AC: Multi-Q with selective forgetting actor-critic for robot path planning"],"prefix":"10.1007","volume":"56","author":[{"ORCID":"https:\/\/orcid.org\/0000-0002-3148-9414","authenticated-orcid":false,"given":"Yuwan","family":"Gu","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Yan","family":"Chen","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Fang","family":"Meng","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Ronghai","family":"Miao","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Jie","family":"Hao","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Jidong","family":"Lv","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2026,2,26]]},"reference":[{"key":"7090_CR1","doi-asserted-by":"publisher","first-page":"34","DOI":"10.1126\/science.153.3731.34","volume":"153","author":"R Bellman","year":"1957","unstructured":"Bellman R (1957) Dynamic programming. Science 153:34\u201337","journal-title":"Science"},{"key":"7090_CR2","unstructured":"Mnih V, Kavukcuoglu K, Silver D, Graves A, Antonoglou I, Wierstra D, Riedmiller MA (2013) Playing Atari with deep reinforcement learning. ArXiv Abs \/1312.5602"},{"key":"7090_CR3","unstructured":"Anschel O, Baram N, Shimkin N (2016) Averaged-dqn: Variance reduction and stabilization for deep reinforcement learning. In: International Conference on Machine Learning https:\/\/api.semanticscholar.org\/CorpusID:2215254"},{"key":"7090_CR4","unstructured":"Silver D, Lever G, Heess N, Degris T, Riedmiller M (2014) Deterministic policy gradient algorithms. PMLR"},{"key":"7090_CR5","doi-asserted-by":"crossref","unstructured":"Navaneethakrishnan M, Pushpa PM, Mohanaprakash TT, Dhanwanth TA (2023) B., S, F.A.A.: Design of biped robot using reinforcement learning and asynchronous actor-critical agent (a3c) algorithm. 2023 2nd International Conference on Vision Towards Emerging Trends in Communication and Networking Technologies (ViTECoN), 1\u20136","DOI":"10.1109\/ViTECoN58111.2023.10156947"},{"key":"7090_CR6","unstructured":"Haarnoja T, Zhou A, Hartikainen K, Tucker G, Ha S, Tan J, Kumar V, Zhu H, Gupta A, Abbeel P, Levine S (2018) Soft actor-critic algorithms and applications. ArXiv abs\/1812.05905"},{"key":"7090_CR7","doi-asserted-by":"publisher","first-page":"47","DOI":"10.1016\/j.ins.2022.08.028","volume":"611","author":"S Chen","year":"2022","unstructured":"Chen S, Qiu X, Tan X, Fang Z, Jin Y (2022) A model-based hybrid soft actor-critic deep reinforcement learning algorithm for optimal ventilator settings. Inf Sci 611:47\u201364. https:\/\/doi.org\/10.1016\/j.ins.2022.08.028","journal-title":"Inf Sci"},{"issue":"4","key":"7090_CR8","doi-asserted-by":"publisher","first-page":"2974","DOI":"10.1109\/TETCI.2024.3369636","volume":"8","author":"X Tan","year":"2024","unstructured":"Tan X, Qu C, Xiong J, Zhang J, Qiu X, Jin Y (2024) Model-based off-policy deep reinforcement learning with model-embedding. IEEE Trans Emerg Top Comput Intell 8(4):2974\u20132986. https:\/\/doi.org\/10.1109\/TETCI.2024.3369636","journal-title":"IEEE Trans Emerg Top Comput Intell"},{"key":"7090_CR9","doi-asserted-by":"crossref","unstructured":"Kainova TD (2023) Overview of the accelerated platform for robotics and artificial intelligence nvidia isaac. 2023 Seminar on Information Computing and Processing (ICP), 89\u201393","DOI":"10.1109\/ICP60417.2023.10397127"},{"key":"7090_CR10","doi-asserted-by":"crossref","unstructured":"Zhang C, Sun P (2023) Heuristic methods for solving the traveling salesman problem (tsp): A comparative study. 2023 IEEE 34th Annual International Symposium on Personal, Indoor and Mobile Radio Communications (PIMRC), 1\u20136","DOI":"10.1109\/PIMRC56721.2023.10293957"},{"key":"7090_CR11","doi-asserted-by":"publisher","unstructured":"Wang L, Weng H, Sun X, Yan W (2023) Research on vehicle routing problem based on heuristic neural network. In: 2023 5th International Conference on Frontiers Technology of Information and Computer (ICFTIC), pp. 1154\u20131158 https:\/\/doi.org\/10.1109\/ICFTIC59930.2023.10456204","DOI":"10.1109\/ICFTIC59930.2023.10456204"},{"key":"7090_CR12","doi-asserted-by":"publisher","first-page":"19111","DOI":"10.1109\/ACCESS.2023.3247730","volume":"11","author":"T Nakamura","year":"2023","unstructured":"Nakamura T, Kobayashi M, Motoi N (2023) Path planning for mobile robot considering turnabouts on narrow road by deep q-network. IEEE Access 11:19111\u201319121","journal-title":"IEEE Access"},{"key":"7090_CR13","doi-asserted-by":"publisher","unstructured":"Chen F, Mei P, Xie H, Yang S, Xu B, Huang C (2022) Reinforcement learningbased energy management control strategy of hybrid electric vehicles. In: 2022 8th International Conference on Control, Automation and Robotics (ICCAR), pp. 248\u2013252 https:\/\/doi.org\/10.1109\/ICCAR55106.2022.9782662","DOI":"10.1109\/ICCAR55106.2022.9782662"},{"issue":"2","key":"7090_CR14","doi-asserted-by":"publisher","DOI":"10.3390\/app10020575","volume":"10","author":"MS Kim","year":"2020","unstructured":"Kim MS, Han DK, Park JH, Kim JS (2020) Motion planning of robot manipulators for a smoother path using a twin delayed deep deterministic policy gradient with hindsight experience replay. Appl Sci 10(2):575","journal-title":"Appl Sci"},{"key":"7090_CR15","doi-asserted-by":"publisher","DOI":"10.1016\/j.oceaneng.2023.115040","author":"X Yang","year":"2023","unstructured":"Yang X, Han Q (2023) Improved reinforcement learning for collision-free local path planning of dynamic obstacle. Ocean Eng. https:\/\/doi.org\/10.1016\/j.oceaneng.2023.115040","journal-title":"Ocean Eng"},{"issue":"1","key":"7090_CR16","doi-asserted-by":"publisher","DOI":"10.1038\/s41598-024-77779-8","volume":"14","author":"L Zanatta","year":"2024","unstructured":"Zanatta L, Barchi F, Manoni S, Tolu S, Bartolini A, Acquaviva A (2024) Exploring spiking neural networks for deep reinforcement learning in robotic tasks. Sci Rep 14(1):30648. https:\/\/doi.org\/10.1038\/s41598-024-77779-8","journal-title":"Sci Rep"},{"key":"7090_CR17","doi-asserted-by":"publisher","unstructured":"Sun Y, Zeng Y, Zhao F, Zhao Z (2023) Multi-compartment neuron and population encoding powered spiking neural network for deep distributional reinforcement learning. https:\/\/doi.org\/10.48550\/arXiv.2301.07275","DOI":"10.48550\/arXiv.2301.07275"},{"key":"7090_CR18","doi-asserted-by":"publisher","unstructured":"Martini M, Eirale A, Cerrato S, Chiaberge M (2023) Pic4rl-gym: a ros2 modular framework for robots autonomous navigation with deep reinforcement learning. In: 2023 3rd International Conference on Computer, Control and Robotics (ICCCR), pp. 198\u2013202 https:\/\/doi.org\/10.1109\/ICCCR56747. 2023.10193996","DOI":"10.1109\/ICCCR56747"},{"issue":"4","key":"7090_CR19","doi-asserted-by":"publisher","first-page":"5064","DOI":"10.1109\/TNNLS.2022.3207346","volume":"35","author":"X Wang","year":"2024","unstructured":"Wang X, Wang S, Liang X, Zhao D, Huang J, Xu X, Dai B, Miao Q (2024) Deep reinforcement learning: a survey. IEEE Trans Neural Networks Learn Syst 35(4):5064\u20135078. https:\/\/doi.org\/10.1109\/TNNLS.2022.3207346","journal-title":"IEEE Trans Neural Networks Learn Syst"},{"key":"7090_CR20","doi-asserted-by":"publisher","DOI":"10.1016\/j.knosys.2025.113689","volume":"322","author":"Y Deng","year":"2025","unstructured":"Deng Y, Qiu X, Chen J, Tan X (2025) Reward guidance for reinforcement learning tasks based on large language models: the lmgt framework. Knowl Based Syst 322:113689. https:\/\/doi.org\/10.1016\/j.knosys.2025.113689","journal-title":"Knowl Based Syst"},{"key":"7090_CR21","doi-asserted-by":"crossref","unstructured":"Carta T, Oudeyer P-Y, Sigaud O, Lamprier S (2022) EAGER: Asking and Answering Questions for Automatic Reward Shaping in Language-guided RL https:\/\/arxiv.org\/abs\/2206.09674","DOI":"10.52202\/068431-0906"},{"key":"7090_CR22","doi-asserted-by":"crossref","unstructured":"Wang H, Tan X, Qiu X, Qu C (2024) Subequivariant reinforcement learning framework for coordinated motion control. 2024 IEEE International Conference on Robotics and Automation (ICRA), 2112\u20132118","DOI":"10.1109\/ICRA57147.2024.10610563"},{"issue":"12","key":"7090_CR23","doi-asserted-by":"publisher","first-page":"8323","DOI":"10.1109\/TAC.2024.3409647","volume":"69","author":"S Meyn","year":"2024","unstructured":"Meyn S (2024) The projected bellman equation in reinforcement learning. IEEE Trans Autom Control 69(12):8323\u20138337. https:\/\/doi.org\/10.1109\/TAC.2024.3409647","journal-title":"IEEE Trans Autom Control"},{"issue":"2","key":"7090_CR24","doi-asserted-by":"publisher","first-page":"317","DOI":"10.1109\/TG.2023.3265975","volume":"16","author":"R-Z Liu","year":"2024","unstructured":"Liu R-Z, Shen Y, Yu Y, Lu T (2024) Revisiting of Alphastar. IEEE Trans Games 16(2):317\u2013330. https:\/\/doi.org\/10.1109\/TG.2023.3265975","journal-title":"IEEE Trans Games"},{"key":"7090_CR25","unstructured":"Gauci J, Conti E, Liang Y, Virochsiri K, He Y, Kaden Z, Narayanan V, Ye X, Fujimoto S (2018) Horizon: Facebook\u2019s open source applied reinforcement learning platform. ArXiv abs\/1811.00260"},{"issue":"6","key":"7090_CR26","doi-asserted-by":"publisher","first-page":"8470","DOI":"10.1109\/TNNLS.2022.3229897","volume":"35","author":"S Disabato","year":"2024","unstructured":"Disabato S, Roveri M (2024) Tiny machine learning for concept drift. IEEE Trans Neural Networks Learn Syst 35(6):8470\u20138481. https:\/\/doi.org\/10.1109\/TNNLS.2022.3229897","journal-title":"IEEE Trans Neural Networks Learn Syst"},{"key":"7090_CR27","unstructured":"Papadimitriou C, Peng B (2023) The complexity of non-stationary reinforcement learning. In: International Conference on Algorithmic Learning Theory https:\/\/api.semanticscholar.org\/CorpusID:259847298"},{"key":"7090_CR28","doi-asserted-by":"publisher","unstructured":"Trat M, Ovtcharova J (2023) Designing concept drift detection ensembles: A survey. In: 2023 IEEE 10th International Conference on Data Science and Advanced Analytics (DSAA), pp. 1\u201310 https:\/\/doi.org\/10.1109\/DSAA60987","DOI":"10.1109\/DSAA60987"},{"key":"7090_CR29","doi-asserted-by":"crossref","unstructured":"Hinder F, Vaquet V, Brinkrolf J, Artelt A, Hammer B (2022) Localization of concept drift: Identifying the drifting datapoints. 2022 International Joint Conference on Neural Networks (IJCNN), 1\u20139","DOI":"10.1109\/IJCNN55064.2022.9892374"},{"issue":"9","key":"7090_CR30","doi-asserted-by":"publisher","first-page":"4243","DOI":"10.1109\/TNNLS.2021.3056201","volume":"33","author":"J Peng","year":"2022","unstructured":"Peng J, Tang B, Jiang H, Li Z, Lei Y, Lin T, Li H (2022) Overcoming long-term catastrophic forgetting through adversarial neural pruning and synaptic consolidation. IEEE Trans Neural Networks Learn Syst 33(9):4243\u20134256. https:\/\/doi.org\/10.1109\/TNNLS.2021.3056201","journal-title":"IEEE Trans Neural Networks Learn Syst"},{"key":"7090_CR31","unstructured":"Lyle C, Rowland M, Dabney W (2022) Understanding and preventing capacity loss in reinforcement learning. ArXiv abs\/2204.09560"},{"key":"7090_CR32","doi-asserted-by":"publisher","unstructured":"Guo R, Yue T (2023) Research on decision control of unmanned driving based on sac. In: 2023 5th International Conference on Frontiers Technology of Information and Computer (ICFTIC), pp. 1104\u20131107 https:\/\/doi.org\/10.1109\/ICFTIC59930.2023.10456323","DOI":"10.1109\/ICFTIC59930.2023.10456323"},{"key":"7090_CR33","doi-asserted-by":"publisher","DOI":"10.1109\/TPAMI.2024.3510627","author":"Shuangming Yang","year":"2025","unstructured":"Yang Shuangming, Linares-Barranco Bernab\u00e9, Wu Yuzhu, Chen BD (2025) Self-supervised high-order information bottleneck learning of spiking neural network for robust event-based optical flow estimation. IEEE Trans Pattern Anal Mach Intell. https:\/\/doi.org\/10.1109\/TPAMI.2024.3510627","journal-title":"IEEE Trans Pattern Anal Mach Intell"},{"issue":"4","key":"7090_CR34","doi-asserted-by":"publisher","first-page":"2156","DOI":"10.1109\/TNNLS.2021.3106044","volume":"34","author":"M Liu","year":"2023","unstructured":"Liu M, Chen L, Du X, Jin L, Shang M (2023) Activated gradients for deep neural networks. IEEE Trans Neural Networks Learn Syst 34(4):2156\u20132168. https:\/\/doi.org\/10.1109\/TNNLS.2021.3106044","journal-title":"IEEE Trans Neural Networks Learn Syst"},{"key":"7090_CR35","unstructured":"Park J, Hwang I, Lee MW, Oh H, Lee MW, Lee Y, Zhang B-T (2022) On the importance of critical period in multi-stage reinforcement learning. ArXiv abs\/2208.04832"},{"key":"7090_CR36","doi-asserted-by":"crossref","unstructured":"Ahn K, Bubeck S, Chewi S, Lee YT, Suarez F, Zhang Y (2023) Learning threshold neurons via edge of stability. In: Neural Inform Process Syst https:\/\/api.semanticscholar.org\/CorpusID:268030620","DOI":"10.52202\/075280-0858"},{"key":"7090_CR37","doi-asserted-by":"publisher","first-page":"4377","DOI":"10.1109\/TIP.2024.3416873","volume":"33","author":"Y Liu","year":"2024","unstructured":"Liu Y, Tian CX, Li H, Wang S (2024) Generalization beyond feature alignment: concept activation-guided contrastive learning. IEEE Trans Image Process 33:4377\u20134390. https:\/\/doi.org\/10.1109\/TIP.2024.3416873","journal-title":"IEEE Trans Image Process"},{"key":"7090_CR38","doi-asserted-by":"publisher","unstructured":"Hwang H, Shin J (2024) Test case prioritization with z-score based neuron coverage. In: 2024 26th International Conference on Advanced Communications Technology (ICACT), pp. 23\u201328 https:\/\/doi.org\/10.23919\/ICACT60172.2024.10471933","DOI":"10.23919\/ICACT60172.2024.10471933"}],"container-title":["Applied Intelligence"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10489-026-07090-8.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s10489-026-07090-8","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10489-026-07090-8.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,3,31]],"date-time":"2026-03-31T05:13:35Z","timestamp":1774934015000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s10489-026-07090-8"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,2]]},"references-count":38,"journal-issue":{"issue":"4","published-print":{"date-parts":[[2026,2]]}},"alternative-id":["7090"],"URL":"https:\/\/doi.org\/10.1007\/s10489-026-07090-8","relation":{},"ISSN":["0924-669X","1573-7497"],"issn-type":[{"value":"0924-669X","type":"print"},{"value":"1573-7497","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,2]]},"assertion":[{"value":"14 October 2024","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"1 January 2026","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"26 February 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Competing interests"}}],"article-number":"117"}}