{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,14]],"date-time":"2026-07-14T22:47:49Z","timestamp":1784069269930,"version":"3.55.0"},"reference-count":41,"publisher":"Springer Science and Business Media LLC","issue":"21","license":[{"start":{"date-parts":[[2024,9,6]],"date-time":"2024-09-06T00:00:00Z","timestamp":1725580800000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2024,9,6]],"date-time":"2024-09-06T00:00:00Z","timestamp":1725580800000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"name":"the Science and Technology Innovation of China","award":["2018AAA0102403"],"award-info":[{"award-number":["2018AAA0102403"]}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Appl Intell"],"published-print":{"date-parts":[[2024,11]]},"DOI":"10.1007\/s10489-024-05771-w","type":"journal-article","created":{"date-parts":[[2024,9,6]],"date-time":"2024-09-06T09:02:24Z","timestamp":1725613344000},"page":"11103-11119","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":2,"title":["Improving multi-UAV cooperative path-finding through multiagent experience learning"],"prefix":"10.1007","volume":"54","author":[{"ORCID":"https:\/\/orcid.org\/0000-0002-4705-1548","authenticated-orcid":false,"given":"Jiang","family":"Longting","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Wei","family":"Ruixuan","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Wang","family":"Dong","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2024,9,6]]},"reference":[{"key":"5771_CR1","doi-asserted-by":"publisher","unstructured":"Bin DH, Feng ZD, Ming FY, Min DY (2019) From wolf pack intelligence to uav swarm cooperative decision-making. Sci Sin Inform 49(112-118). https:\/\/doi.org\/10.1360\/N112018-00168","DOI":"10.1360\/N112018-00168"},{"key":"5771_CR2","unstructured":"Yu\u00a0YP, Duan HB Yuan WM (2022) Pursuit-evasion control for uav swarm imitating the intelligent behavior in hawks-starlings. J Command Control 8(422-433)"},{"key":"5771_CR3","doi-asserted-by":"publisher","unstructured":"Feng L, Ruixuan W, Kai Z, Chao D (2022) Research on multi-uav roundup strategy based on the unity of group will. J Beijing University Aeronaut Astronaut 48(2241-2249). https:\/\/doi.org\/10.13700\/j.bh.1001-5965.2021.0109","DOI":"10.13700\/j.bh.1001-5965.2021.0109"},{"issue":"11","key":"5771_CR4","doi-asserted-by":"publisher","first-page":"4948","DOI":"10.3390\/app11114948","volume":"11","author":"L Canese","year":"2021","unstructured":"Canese L, Cardarilli GC, Di Nunzio L, Fazzolari R, Giardino D, Re M, Span\u00f2 S (2021) Multi-agent reinforcement learning: A review of challenges and applications. Appl Sci 11(11):4948. https:\/\/doi.org\/10.3390\/app11114948","journal-title":"Appl Sci"},{"issue":"5","key":"5771_CR5","doi-asserted-by":"publisher","first-page":"3215","DOI":"10.1007\/s10462-020-09938-y","volume":"54","author":"W Du","year":"2021","unstructured":"Du W, Ding S (2021) A survey on multi-agent deep reinforcement learning: from the perspective of challenges and applications. Artif Intell Rev 54(5):3215\u20133238. https:\/\/doi.org\/10.1007\/s10462-020-09938-y","journal-title":"Artif Intell Rev"},{"issue":"3","key":"5771_CR6","doi-asserted-by":"publisher","first-page":"5175","DOI":"10.1007\/s10586-017-1132-9","volume":"22","author":"Y Cao","year":"2019","unstructured":"Cao Y, Wei W, Bai Y, Qiao H (2019) Multi-base multi-uav cooperative reconnaissance path planning with genetic algorithm. Clust Comput J Netw Softw Tools Appl 22(3):5175\u20135184. https:\/\/doi.org\/10.1007\/s10586-017-1132-9","journal-title":"Clust Comput J Netw Softw Tools Appl"},{"issue":"07","key":"5771_CR7","doi-asserted-by":"publisher","first-page":"1551","DOI":"10.3969\/j.issn.1001-506X.2019.07.16","volume":"41","author":"T Hu","year":"2019","unstructured":"Hu T, Liu ZJ, Liu Y, Xia SS, Chen QB (2019) Multi-uav 3d reconnaissance path planning. Phys. Rev. E. 41(07):1551\u20131559. https:\/\/doi.org\/10.3969\/j.issn.1001-506X.2019.07.16","journal-title":"Phys. Rev. E."},{"issue":"02","key":"5771_CR8","doi-asserted-by":"publisher","first-page":"137","DOI":"10.3969\/j.issn.1673-3193.2019.02.008","volume":"40","author":"JC Niu","year":"2019","unstructured":"Niu JC (2019) Zhang PJ Wang Z Q: Path planning based on optimal ant colony algorithm in multi-machine cooperative operation. J North China Univ (Nat Sci Ed) 40(02):137\u2013142. https:\/\/doi.org\/10.3969\/j.issn.1673-3193.2019.02.008","journal-title":"J North China Univ (Nat Sci Ed)"},{"key":"5771_CR9","unstructured":"Jin L, Liu GX, Hui ZJ (2024) Particle swarm optimization algorithm based on labor division and fuzzy control. Complex Syst Complex Sci"},{"key":"5771_CR10","doi-asserted-by":"publisher","unstructured":"Gang XD, Xin GY Qing WZ (2023) Review of whale optimization algorithm. Appl Res Comput 40(328-336). https:\/\/doi.org\/10.19734\/j.issn.1001-3695.2022.06.0347","DOI":"10.19734\/j.issn.1001-3695.2022.06.0347"},{"key":"5771_CR11","doi-asserted-by":"publisher","first-page":"849","DOI":"10.1016\/j.future.2019.02.028","volume":"97","author":"AA Heidari","year":"2019","unstructured":"Heidari AA, Mirjalili S, Faris H, Aljarah I, Mafarja M, Chen H (2019) Harris hawks optimization: Algorithm and applications. Futur Gener Comput Syst Int J Escience 97:849\u2013872. https:\/\/doi.org\/10.1016\/j.future.2019.02.028","journal-title":"Futur Gener Comput Syst Int J Escience"},{"issue":"6","key":"5771_CR12","doi-asserted-by":"publisher","first-page":"2201","DOI":"10.1007\/s10489-018-1384-y","volume":"49","author":"RK Dewangan","year":"2019","unstructured":"Dewangan RK, Shukla A, Godfrey WW (2019) Three dimensional path planning using grey wolf optimizer for uavs. Appl Intell 49(6):2201\u20132217. https:\/\/doi.org\/10.1007\/s10489-018-1384-y","journal-title":"Appl Intell"},{"key":"5771_CR13","doi-asserted-by":"publisher","first-page":"196","DOI":"10.1016\/j.comcom.2020.04.050","volume":"162","author":"C Xu","year":"2020","unstructured":"Xu C, Xu M, Yin C (2020) Optimized multi-uav cooperative path planning under the complex confrontation environment. Comput Commun 162:196\u2013203. https:\/\/doi.org\/10.1016\/j.comcom.2020.04.050","journal-title":"Comput Commun"},{"key":"5771_CR14","doi-asserted-by":"publisher","first-page":"229","DOI":"10.1016\/j.neucom.2018.06.032","volume":"313","author":"D Zhang","year":"2018","unstructured":"Zhang D, Duan H (2018) Social-class pigeon-inspired optimization and time stamp segmentation for multi-uav cooperative path planning. Neurocomputing 313:229\u2013246. https:\/\/doi.org\/10.1016\/j.neucom.2018.06.032","journal-title":"Neurocomputing"},{"key":"5771_CR15","doi-asserted-by":"publisher","first-page":"146264","DOI":"10.1109\/ACCESS.2019.2943253","volume":"7","author":"H Qie","year":"2019","unstructured":"Qie H, Shi D, Shen T, Xu X, Li Y, Wang L (2019) Joint optimization of multi-uav target assignment and path planning based on multi-agent reinforcement learning. IEEE ACCESS 7:146264\u2013146272. https:\/\/doi.org\/10.1109\/ACCESS.2019.2943253","journal-title":"IEEE ACCESS"},{"key":"5771_CR16","doi-asserted-by":"publisher","first-page":"410","DOI":"10.1016\/j.neucom.2020.06.038","volume":"410","author":"X Lan","year":"2020","unstructured":"Lan X, Liu Y, Zhao Z (2020) Cooperative control for swarming systems based on reinforcement learning in unknown dynamic environment. Neurocomputing 410:410\u2013418. https:\/\/doi.org\/10.1016\/j.neucom.2020.06.038","journal-title":"Neurocomputing"},{"issue":"7","key":"5771_CR17","doi-asserted-by":"publisher","first-page":"100","DOI":"10.1016\/j.cja.2021.09.008","volume":"35","author":"Z Wenhong","year":"2022","unstructured":"Wenhong Z, Jie L, Zhihong L, Lincheng S (2022) Improving multi-target cooperative tracking guidance for uav swarms using multi-agent reinforcement learning. Chin J Aeronaut 35(7):100\u2013112. https:\/\/doi.org\/10.1016\/j.cja.2021.09.008","journal-title":"Chin J Aeronaut"},{"key":"5771_CR18","doi-asserted-by":"publisher","unstructured":"Jiang L, Wei R, Wang D (2022) Uavs rounding up inspired by communication multi-agent depth deterministic policy gradient. Appl Intell 1\u201316. https:\/\/doi.org\/10.1007\/s10489-022-03986-3","DOI":"10.1007\/s10489-022-03986-3"},{"key":"5771_CR19","doi-asserted-by":"publisher","unstructured":"Gao J, Shi X, Yu JJQ (2021) Attn-commnet: Coordinated traffic lights control on large-scale network level. In: 2021 IEEE 33rd International conference on tools with artificial intelligence (ICTAI), pp 289\u2013293. https:\/\/doi.org\/10.1109\/ICTAI52525.2021.00048","DOI":"10.1109\/ICTAI52525.2021.00048"},{"key":"5771_CR20","doi-asserted-by":"publisher","first-page":"4195","DOI":"10.1007\/s10489-020-01755-8","volume":"50","author":"H Chen","year":"2020","unstructured":"Chen H, Liu Y, Zhou Z, Hu D, Zhang M (2020) Gama: Graph attention multi-agent reinforcement learning algorithm for cooperation. Appl Intell 50:4195\u20134205. https:\/\/doi.org\/10.1007\/s10489-020-01755-8","journal-title":"Appl Intell"},{"issue":"1","key":"5771_CR21","first-page":"7234","volume":"21","author":"T Rashid","year":"2020","unstructured":"Rashid T, Samvelyan M, De Witt CS, Farquhar G, Foerster J, Whiteson S (2020) Monotonic value function factorisation for deep multi-agent reinforcement learning. J Mach Learn Res 21(1):7234\u20137284","journal-title":"J Mach Learn Res"},{"key":"5771_CR22","unstructured":"Mahajan A, Rashid T, Samvelyan M, Whiteson S (2019) Maven: Multi-agent variational exploration. Adv Neural Inf Process Syst 32"},{"key":"5771_CR23","doi-asserted-by":"publisher","unstructured":"Huang L, Fu M, Rao A, Irissappane AA, Zhang J, Xu C (2022) A distributional perspective on multiagent cooperation with deep reinforcement learning. IEEE Trans Neural Netw Learn Syst 1\u201314. https:\/\/doi.org\/10.1109\/TNNLS.2022.3202097","DOI":"10.1109\/TNNLS.2022.3202097"},{"issue":"2","key":"5771_CR24","doi-asserted-by":"publisher","first-page":"82","DOI":"10.1504\/IJBIC.2021.118087","volume":"18","author":"B Li","year":"2021","unstructured":"Li B, Liang S, Gan Z, Chen D, Gao P (2021) Research on multi-uav task decision-making based on improved maddpg algorithm and transfer learning. Int J Bio-Inspired Comput 18(2):82\u201391. https:\/\/doi.org\/10.1504\/IJBIC.2021.118087","journal-title":"Int J Bio-Inspired Comput"},{"issue":"12","key":"5771_CR25","doi-asserted-by":"publisher","first-page":"10497","DOI":"10.1109\/JIOT.2023.3240173","volume":"10","author":"H Kang","year":"2023","unstructured":"Kang H, Chang X, Mi\u0161i\u0107 J, Mi\u0161i\u0107 VB, Fan J, Liu Y (2023) Cooperative uav resource allocation and task offloading in hierarchical aerial computing systems: A mappo-based approach. IEEE Internet Things J 10(12):10497\u201310509. https:\/\/doi.org\/10.1109\/JIOT.2023.3240173","journal-title":"IEEE Internet Things J"},{"key":"5771_CR26","doi-asserted-by":"publisher","unstructured":"Liu X, Yin Y, Su Y, Ming R (2022) A multi-ucav cooperative decision-making method based on an mappo algorithm for beyond-visual-range air combat. Aerospace 9(10). https:\/\/doi.org\/10.3390\/aerospace9100563","DOI":"10.3390\/aerospace9100563"},{"issue":"11","key":"5771_CR27","doi-asserted-by":"publisher","first-page":"12597","DOI":"10.1109\/TVT.2020.3026111","volume":"69","author":"Y Guan","year":"2020","unstructured":"Guan Y, Ren Y, Li SE, Sun Q, Luo L, Li K (2020) Centralized cooperation for connected and automated vehicles at intersections by proximal policy optimization. IEEE Trans Veh Technol 69(11):12597\u201312608. https:\/\/doi.org\/10.1109\/TVT.2020.3026111","journal-title":"IEEE Trans Veh Technol"},{"issue":"2","key":"5771_CR28","doi-asserted-by":"publisher","first-page":"502","DOI":"10.1007\/s10489-019-01527-z","volume":"50","author":"F Shoeleh","year":"2020","unstructured":"Shoeleh F, Asadpour M (2020) Skill based transfer learning with domain adaptation for continuous reinforcement learning domains. Appl Intell 50(2):502\u2013518. https:\/\/doi.org\/10.1007\/s10489-019-01527-z","journal-title":"Appl Intell"},{"key":"5771_CR29","unstructured":"Ji ZX (2016) Research on adaptive recommendation algorithm based on experience learning. Master\u2019s thesis, Dalian University of Technologys"},{"key":"5771_CR30","unstructured":"Wang H (2019) Research on map construction and path planning technology of mobile robot in indoor environment. Master\u2019s thesis, An-Hui Engineering University"},{"key":"5771_CR31","unstructured":"Shi Y J (2017) Improvement of swarm intelligence algorithm and its application analysis. Master\u2019s thesis, Nanjing University of Posts and Telecommunications. https:\/\/doi.org\/CNKI:CDMD:2.1017.859356"},{"issue":"1","key":"5771_CR32","doi-asserted-by":"publisher","first-page":"35","DOI":"10.1007\/s12369-019-00531-0","volume":"12","author":"R Wei","year":"2020","unstructured":"Wei R, Zhang Q, Xu Z (2020) Peers\u2019 experience learning for developmental robots. Int J Soc Robot 12(1):35\u201345. https:\/\/doi.org\/10.1007\/s12369-019-00531-0","journal-title":"Int J Soc Robot"},{"issue":"S2","key":"5771_CR33","doi-asserted-by":"publisher","DOI":"10.7527\/S1000-6893.2020.24285","volume":"40","author":"K Zhou","year":"2020","unstructured":"Zhou K, Wei R, Zhang Q, Ding C (2020) Learning method for autonomous air combat based on experience transfer. Acta Aeronautica et Astronautica Sinica 40(S2):724285. https:\/\/doi.org\/10.7527\/S1000-6893.2020.24285","journal-title":"Acta Aeronautica et Astronautica Sinica"},{"key":"5771_CR34","doi-asserted-by":"publisher","first-page":"8129","DOI":"10.1109\/ACCESS.2020.2964031","volume":"8","author":"K Zhou","year":"2020","unstructured":"Zhou K, Wei R, Zhang Q, Xu Z (2020) Learning system for air combat decision inspired by cognitive mechanisms of the brain. IEEE Access 8:8129\u20138144. https:\/\/doi.org\/10.1109\/ACCESS.2020.2964031","journal-title":"IEEE Access"},{"key":"5771_CR35","doi-asserted-by":"publisher","unstructured":"B\u00f8hn E, Coates EM, Moe S, Johansen TA (2019) Deep reinforcement learning attitude control of fixed-wing uavs using proximal policy optimization. In: 2019 International conference on unmanned aircraft systems (ICUAS), pp 523\u2013533. https:\/\/doi.org\/10.1109\/ICUAS.2019.8798254. IEEE","DOI":"10.1109\/ICUAS.2019.8798254"},{"key":"5771_CR36","first-page":"13458","volume":"34","author":"JG Kuba","year":"2021","unstructured":"Kuba JG, Wen M, Meng L, Zhang H, Mguni D, Wang J, Yang Y et al (2021) Settling the variance of multi-agent policy gradients. Adv Neural Inf Process Syst 34:13458\u201313470","journal-title":"Adv Neural Inf Process Syst"},{"key":"5771_CR37","unstructured":"Kuba JG, Chen R, Wen M, Wen Y, Sun F, Wang J, Yang Y (2021) Trust region policy optimisation in multi-agent reinforcement learning. arXiv preprint arXiv:2109.11251"},{"issue":"2","key":"5771_CR38","doi-asserted-by":"publisher","first-page":"249","DOI":"10.1109\/JAS.2021.1003814","volume":"8","author":"D Bertsekas","year":"2021","unstructured":"Bertsekas D (2021) Multiagent reinforcement learning: Rollout and policy iteration. IEEE-CAA J Autom Sin 8(2):249\u2013272. https:\/\/doi.org\/10.1109\/JAS.2021.1003814","journal-title":"IEEE-CAA J Autom Sin"},{"key":"5771_CR39","doi-asserted-by":"publisher","unstructured":"Schulman J, Moritz P, Levine S, Jordan M, Abbeel P (2015) High-dimensional continuous control using generalized advantage estimation. arXiv preprint arXiv:1506.02438. https:\/\/doi.org\/10.48550\/arXiv.1506.02438","DOI":"10.48550\/arXiv.1506.02438"},{"issue":"1","key":"5771_CR40","first-page":"954","volume":"14","author":"E Jacinto","year":"2023","unstructured":"Jacinto E, Martinez F, Martinez F (2023) Navigation of autonomous vehicles using reinforcement learning with generalized advantage estimation. Int J Adv Comput Sci Appl 14(1):954\u2013959","journal-title":"Int J Adv Comput Sci Appl"},{"key":"5771_CR41","doi-asserted-by":"publisher","first-page":"100017","DOI":"10.1016\/j.commtr.2021.100017","volume":"1","author":"B Peng","year":"2021","unstructured":"Peng B, Keskin MF, Kulcs\u00e1r B, Wymeersch H (2021) Connected autonomous vehicles for improving mixed traffic efficiency in unsignalized intersections with deep reinforcement learning. Commun Transp Res 1:100017. https:\/\/doi.org\/10.1016\/j.commtr.2021.100017","journal-title":"Commun Transp Res"}],"container-title":["Applied Intelligence"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10489-024-05771-w.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s10489-024-05771-w\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10489-024-05771-w.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2024,9,18]],"date-time":"2024-09-18T14:32:54Z","timestamp":1726669974000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s10489-024-05771-w"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2024,9,6]]},"references-count":41,"journal-issue":{"issue":"21","published-print":{"date-parts":[[2024,11]]}},"alternative-id":["5771"],"URL":"https:\/\/doi.org\/10.1007\/s10489-024-05771-w","relation":{},"ISSN":["0924-669X","1573-7497"],"issn-type":[{"value":"0924-669X","type":"print"},{"value":"1573-7497","type":"electronic"}],"subject":[],"published":{"date-parts":[[2024,9,6]]},"assertion":[{"value":"12 August 2024","order":1,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"6 September 2024","order":2,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Competing interest"}}]}}