{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,2]],"date-time":"2026-07-02T08:59:28Z","timestamp":1782982768072,"version":"3.54.5"},"reference-count":44,"publisher":"Springer Science and Business Media LLC","issue":"4","license":[{"start":{"date-parts":[[2024,10,10]],"date-time":"2024-10-10T00:00:00Z","timestamp":1728518400000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2024,10,10]],"date-time":"2024-10-10T00:00:00Z","timestamp":1728518400000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100001809","name":"the National Natural Science Foundation of China","doi-asserted-by":"crossref","award":["61873086"],"award-info":[{"award-number":["61873086"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"crossref"}]},{"DOI":"10.13039\/501100013058","name":"Jiangsu Provincial Key Research and Development Program","doi-asserted-by":"publisher","award":["BE2023340"],"award-info":[{"award-number":["BE2023340"]}],"id":[{"id":"10.13039\/501100013058","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Int. J. Mach. Learn. &amp; Cyber."],"published-print":{"date-parts":[[2025,4]]},"DOI":"10.1007\/s13042-024-02393-z","type":"journal-article","created":{"date-parts":[[2024,10,10]],"date-time":"2024-10-10T03:41:06Z","timestamp":1728531666000},"page":"2315-2333","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":9,"title":["Improved D3QN with graph augmentation for enhanced multi-UAV cooperative path planning in urban environments"],"prefix":"10.1007","volume":"16","author":[{"given":"Yonghao","family":"Zhao","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Jianjun","family":"Ni","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Guangyi","family":"Tang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Yang","family":"Gu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Simon X.","family":"Yang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2024,10,10]]},"reference":[{"key":"2393_CR1","doi-asserted-by":"crossref","unstructured":"Aslan S (2024) A hospitalization mechanism based immune plasma algorithm for path planning of unmanned aerial vehicles. International Journal of Machine Learning and Cybernetics (2024, Article in Press)","DOI":"10.1007\/s13042-023-02087-y"},{"key":"2393_CR2","doi-asserted-by":"publisher","first-page":"2465","DOI":"10.3390\/rs16132465","volume":"16","author":"J Ni","year":"2024","unstructured":"Ni J, Zhu S, Tang G, Ke C, Wang T (2024) A small-object detection model based on improved yolov8s for uav image scenarios. Remote Sensing 16:2465","journal-title":"Remote Sensing"},{"issue":"1","key":"2393_CR3","doi-asserted-by":"publisher","first-page":"76","DOI":"10.20517\/ir.2023.04","volume":"3","author":"Z Zheng","year":"2023","unstructured":"Zheng Z, Duan H (2023) Uav maneuver decision-making via deep reinforcement learning for short-range air combat. Intelligence & Robotics 3(1):76\u201394","journal-title":"Intelligence & Robotics"},{"issue":"5","key":"2393_CR4","doi-asserted-by":"publisher","first-page":"4933","DOI":"10.1109\/TIE.2023.3285921","volume":"71","author":"Y Zhao","year":"2024","unstructured":"Zhao Y, Yan L, Xie H, Dai J, Wei P (2024) Autonomous exploration method for fast unknown environment mapping by using uav equipped with limited fov sensor. IEEE Trans Industr Electron 71(5):4933\u20134943","journal-title":"IEEE Trans Industr Electron"},{"issue":"19","key":"2393_CR5","doi-asserted-by":"publisher","first-page":"4954","DOI":"10.3390\/rs14194954","volume":"14","author":"A Lambertini","year":"2022","unstructured":"Lambertini A, Mandanici E, Tini MA, Vittuari L (2022) Technical challenges for multi-temporal and multi-sensor image processing surveyed by uav for mapping and monitoring in precision agriculture. Remote Sensing 14(19):4954","journal-title":"Remote Sensing"},{"key":"2393_CR6","doi-asserted-by":"publisher","first-page":"222","DOI":"10.1016\/j.isatra.2023.01.007","volume":"137","author":"Y Wang","year":"2023","unstructured":"Wang Y, Liu W, Liu J, Sun C (2023) Cooperative usv-uav marine search and rescue with visual navigation and reinforcement learning-based control. ISA Trans 137:222\u2013235","journal-title":"ISA Trans"},{"key":"2393_CR7","doi-asserted-by":"publisher","DOI":"10.1016\/j.eswa.2023.121495","volume":"237","author":"G Paulin","year":"2024","unstructured":"Paulin G, Sambolek S, Ivasic-Kos M (2024) Application of raycast method for person geolocalization and distance determination using uav images in real-world land search and rescue scenarios. Expert Syst Appl 237:121495","journal-title":"Expert Syst Appl"},{"issue":"4","key":"2393_CR8","doi-asserted-by":"publisher","first-page":"1537","DOI":"10.1109\/TSMC.2018.2815988","volume":"50","author":"HX Pham","year":"2020","unstructured":"Pham HX, La HM, Feil-Seifer D, Deans MC (2020) A distributed control framework of multiple unmanned aerial vehicles for dynamic wildfire tracking. IEEE Transactions on Systems, Man, and Cybernetics: Systems 50(4):1537\u20131548","journal-title":"IEEE Transactions on Systems, Man, and Cybernetics: Systems"},{"key":"2393_CR9","doi-asserted-by":"crossref","unstructured":"De\u00a0Lima\u00a0Filho G.M, Kuroswiski A.R, Medeiros F.L.L, Voskuijl M, Monsuur H, Passaro A (2022). Optimization of unmanned air vehicle tactical formation in war games. IEEE Access 10, 21727\u201321741","DOI":"10.1109\/ACCESS.2022.3152768"},{"key":"2393_CR10","doi-asserted-by":"crossref","unstructured":"Zhang Y, Zhao W, Wang J, Yuan Y (2024). Recent progress, challenges and future prospects of applied deep reinforcement learning: A practical perspective in path planning. Neurocomputing 608","DOI":"10.1016\/j.neucom.2024.128423"},{"key":"2393_CR11","doi-asserted-by":"crossref","unstructured":"Ganesan S, Ramalingam B, Mohan R.E (2024). A hybrid sampling-based rrt* path planning algorithm for autonomous mobile robot navigation. Expert Systems with Applications 258","DOI":"10.1016\/j.eswa.2024.125206"},{"key":"2393_CR12","doi-asserted-by":"crossref","unstructured":"Javed S, Hassan A, Ahmad R, Ahmed W, Ahmed R, Saadat A, Guizani M (2024). State-of-the-art and future research challenges in uav swarms. IEEE Internet of Things Journal, 1\u20131","DOI":"10.1109\/JIOT.2024.3364230"},{"key":"2393_CR13","doi-asserted-by":"crossref","unstructured":"Liu J, Liao X, Ye H, Yue H, Wang Y, Tan X, Wang D (2022). Uav swarm scheduling method for remote sensing observations during emergency scenarios. Remote Sensing 14(6)","DOI":"10.3390\/rs14061406"},{"key":"2393_CR14","doi-asserted-by":"publisher","first-page":"196","DOI":"10.1016\/j.comcom.2020.04.050","volume":"162","author":"C Xu","year":"2020","unstructured":"Xu C, Xu M, Yin C (2020) Optimized multi-uav cooperative path planning under the complex confrontation environment. Comput Commun 162:196\u2013203","journal-title":"Comput Commun"},{"issue":"1","key":"2393_CR15","doi-asserted-by":"publisher","first-page":"123","DOI":"10.1109\/TCYB.2023.3265926","volume":"54","author":"J Fu","year":"2024","unstructured":"Fu J, Sun G, Liu J, Yao W, Wu L (2024) On hierarchical multi-uav dubins traveling salesman problem paths in a complex obstacle environment. IEEE Transactions on Cybernetics 54(1):123\u2013135","journal-title":"IEEE Transactions on Cybernetics"},{"issue":"5","key":"2393_CR16","doi-asserted-by":"publisher","first-page":"5723","DOI":"10.1109\/TVT.2020.2982508","volume":"69","author":"Q Liu","year":"2020","unstructured":"Liu Q, Shi L, Sun L, Li J, Ding M, Shu FS (2020) Path planning for uav-mounted mobile edge computing with deep reinforcement learning. IEEE Trans Veh Technol 69(5):5723\u20135728","journal-title":"IEEE Trans Veh Technol"},{"key":"2393_CR17","doi-asserted-by":"crossref","unstructured":"Silvirianti Narottama B, Shin S.Y (2023). Uav coverage path planning with quantum-based recurrent deep deterministic policy gradient. IEEE Transactions on Vehicular Technology, 1\u20136","DOI":"10.36227\/techrxiv.21973784.v1"},{"issue":"3","key":"2393_CR18","doi-asserted-by":"publisher","first-page":"374","DOI":"10.20517\/ir.2023.22","volume":"3","author":"J Ni","year":"2023","unstructured":"Ni J, Chen Y, Tang G, Shi J, Cao WC, Shi P (2023) Deep learning-based scene understanding for autonomous robots: a survey. Intelligence & Robotics 3(3):374\u2013401","journal-title":"Intelligence & Robotics"},{"key":"2393_CR19","doi-asserted-by":"crossref","unstructured":"Gao Z, Zhang X, Li Y, Zhu Y, Wu H, Guan X (2022). Analyses and comparisons of uav path planning algorithms in three-dimensional city environment. In: IEEE Conference on Intelligent Transportation Systems, Proceedings, ITSC, Macau, China, 459\u2013464","DOI":"10.1109\/ITSC55140.2022.9922063"},{"key":"2393_CR20","doi-asserted-by":"crossref","unstructured":"Zhou Q, Liu G (2022). Uav path planning based on the combination of a-star algorithm and rrt-star algorithm. In: Proceedings of 2022 IEEE International Conference on Unmanned Systems, ICUS 2022, Guangzhou, China, 146\u2013151","DOI":"10.1109\/ICUS55513.2022.9986703"},{"key":"2393_CR21","doi-asserted-by":"crossref","unstructured":"Yu Z, Chen Y (2023). Persistent monitoring uav path planning based on entropy optimization. In: Proceedings of 13th IEEE International Conference on CYBER Technology in Automation, Control, and Intelligent Systems, CYBER 2023, Qinhuangdao, China, 909\u2013914","DOI":"10.1109\/CYBER59472.2023.10256557"},{"issue":"1","key":"2393_CR22","doi-asserted-by":"publisher","first-page":"231","DOI":"10.1007\/s12555-021-0666-z","volume":"21","author":"D Seo","year":"2023","unstructured":"Seo D, Kang J (2023) Collision-avoided tracking control of uav using velocity-adaptive 3d local path planning. Int J Control Autom Syst 21(1):231\u2013243","journal-title":"Int J Control Autom Syst"},{"key":"2393_CR23","doi-asserted-by":"crossref","unstructured":"Wang Z, Wan C, Lv X, Ni C, Mao Z, Li Y (2023). Multi-uav online path planning algorithm based on improved hybrid a. In: 2023 6th International Symposium on Autonomous Systems, ISAS 2023, Nanjing, China, 1\u20136","DOI":"10.1109\/ISAS59543.2023.10164537"},{"key":"2393_CR24","doi-asserted-by":"crossref","unstructured":"Huang H, Li H, Wang M, Wu Y, He X (2022). Multi-uav cooperative path planning based on aquila optimizer. In: International Conference on Autonomous Unmanned Systems, Xi\u2019an, China, 2005\u20132014","DOI":"10.1007\/978-981-99-0479-2_186"},{"key":"2393_CR25","doi-asserted-by":"crossref","unstructured":"Ma Y.K, Li S.R (2023). Uav path planning based on improved artificial potential field method. In: Lecture Notes in Electrical Engineering, Ningbo, China, 761\u2013777","DOI":"10.1007\/978-981-99-6882-4_62"},{"key":"2393_CR26","doi-asserted-by":"crossref","unstructured":"Zhang Z, Liu S, Zhou J, Yin Y, Jia H, Ma L (2021). Survey of uav path planning based on swarm intelligence optimization. In: 10th International Conference on Communications, Signal Processing, and Systems, CSPS 2021, Changbaishan, China, 318\u2013326","DOI":"10.1007\/978-981-19-0390-8_39"},{"issue":"5","key":"2393_CR27","doi-asserted-by":"publisher","first-page":"4295","DOI":"10.1007\/s10462-022-10281-7","volume":"56","author":"J Tang","year":"2023","unstructured":"Tang J, Duan H, Lao S (2023) Swarm intelligence algorithms for multiple unmanned aerial vehicles collaboration: a comprehensive review. Artif Intell Rev 56(5):4295\u20134327","journal-title":"Artif Intell Rev"},{"issue":"17","key":"2393_CR28","doi-asserted-by":"publisher","first-page":"19881","DOI":"10.1109\/JSEN.2023.3297666","volume":"23","author":"L Zu","year":"2023","unstructured":"Zu L, Wang Z, Liu C, Ge SS (2023) Research on uav path planning method based on improved hpo algorithm in multitask environment. IEEE Sens J 23(17):19881\u201319893","journal-title":"IEEE Sens J"},{"key":"2393_CR29","doi-asserted-by":"crossref","unstructured":"Gao H, Bai H (2023). Uav path planning method based on quantum squirrel search algorithm. In: 2023 IEEE International Conference on Mechatronics and Automation, ICMA 2023, Harbin, Heilongjiang, China, 1883\u20131887","DOI":"10.1109\/ICMA57826.2023.10215557"},{"issue":"4","key":"2393_CR30","doi-asserted-by":"publisher","first-page":"2658","DOI":"10.1109\/TCYB.2022.3170580","volume":"53","author":"Y Wan","year":"2023","unstructured":"Wan Y, Zhong Y, Ma A, Zhang L (2023) An accurate uav 3-d path planning method for disaster emergency response based on an improved multiobjective swarm intelligence algorithm. IEEE Transactions on Cybernetics 53(4):2658\u20132671","journal-title":"IEEE Transactions on Cybernetics"},{"issue":"3","key":"2393_CR31","doi-asserted-by":"publisher","first-page":"1659","DOI":"10.1109\/COMST.2021.3073036","volume":"23","author":"W Chen","year":"2021","unstructured":"Chen W, Qiu X, Cai T, Dai H-N, Zheng Z, Zhang Y (2021) Deep reinforcement learning for internet of things: A comprehensive survey. IEEE Communications Surveys and Tutorials 23(3):1659\u20131692","journal-title":"IEEE Communications Surveys and Tutorials"},{"key":"2393_CR32","unstructured":"Nikpour B, Sinodinos D, Armanfard N (2024). Deep reinforcement learning in human activity recognition: A survey and outlook. IEEE Transactions on Neural Networks and Learning Systems, 1\u201312"},{"issue":"10","key":"2393_CR33","doi-asserted-by":"publisher","first-page":"9725","DOI":"10.1109\/TVT.2021.3102589","volume":"70","author":"D Hong","year":"2021","unstructured":"Hong D, Lee S, Cho YH, Baek D, Kim J, Chang N (2021) Energy-efficient online path planning of multiple drones using reinforcement learning. IEEE Trans Veh Technol 70(10):9725\u20139740","journal-title":"IEEE Trans Veh Technol"},{"issue":"12","key":"2393_CR34","doi-asserted-by":"publisher","first-page":"13022","DOI":"10.1109\/TVT.2021.3121747","volume":"70","author":"H Xie","year":"2021","unstructured":"Xie H, Yang D, Xiao L, Lyu J (2021) Connectivity-aware 3d uav path design with deep reinforcement learning. IEEE Trans Veh Technol 70(12):13022\u201313034","journal-title":"IEEE Trans Veh Technol"},{"issue":"7","key":"2393_CR35","doi-asserted-by":"publisher","first-page":"5649","DOI":"10.1007\/s00521-021-06702-3","volume":"34","author":"MN Alpdemir","year":"2022","unstructured":"Alpdemir MN (2022) Tactical uav path optimization under radar threat using deep reinforcement learning. Neural Comput Appl 34(7):5649\u20135664","journal-title":"Neural Comput Appl"},{"key":"2393_CR36","doi-asserted-by":"crossref","unstructured":"Wu J, Sun Y, Li D, Shi J, Li X, Gao L, Yu L, Han G, Wu J (2023). An adaptive conversion speed q-learning algorithm for search and rescue uav path planning in unknown environments. IEEE Transactions on Vehicular Technology, 1\u201314","DOI":"10.1109\/TVT.2023.3297837"},{"key":"2393_CR37","doi-asserted-by":"publisher","first-page":"18","DOI":"10.3390\/drones8010018","volume":"8","author":"X Zhao","year":"2024","unstructured":"Zhao X, Yang R, Zhong L, Hou Z (2024) Multi-uav path planning and following based on multi-agent reinforcement learning. Drones 8:18","journal-title":"Drones"},{"key":"2393_CR38","doi-asserted-by":"publisher","first-page":"1302898","DOI":"10.3389\/fnbot.2023.1302898","volume":"17","author":"X Kong","year":"2023","unstructured":"Kong X, Zhou Y, Li Z, Wang S (2023) Multi-uav simultaneous target assignment and path planning based on deep reinforcement learning in dynamic multiple obstacles environments. Front Neurorobot 17:1302898","journal-title":"Front Neurorobot"},{"key":"2393_CR39","doi-asserted-by":"crossref","unstructured":"Mnih V, Kavukcuoglu K, Silver D, Rusu A.A, Veness J, Bellemare M.G, Graves A, Riedmiller M, Fidjeland A.K, Ostrovski G, et al (2015). Human-level control through deep reinforcement learning. nature 518(7540), 529\u2013533","DOI":"10.1038\/nature14236"},{"issue":"4","key":"2393_CR40","doi-asserted-by":"publisher","first-page":"1031","DOI":"10.1007\/s13042-020-01218-z","volume":"12","author":"PPK Chan","year":"2021","unstructured":"Chan PPK, Xiao M, Qin X, Kees N (2021) Dynamic fusion for ensemble of deep q-network. Int J Mach Learn Cybern 12(4):1031\u20131040","journal-title":"Int J Mach Learn Cybern"},{"key":"2393_CR41","unstructured":"Wang Z, Schaul T, Hessel M, Van\u00a0Hasselt H, Lanctot M, De\u00a0Frcitas N (2016).Dueling network architectures for deep reinforcement learning. In: 33rd International Conference on Machine Learning, ICML 2016, vol. 4. New York City, NY, United states, 2939\u20132947"},{"key":"2393_CR42","doi-asserted-by":"crossref","unstructured":"Van\u00a0Hasselt H, Guez A, Silver D (2016). Deep reinforcement learning with double q-learning. In: 30th AAAI Conference on Artificial Intelligence, AAAI 2016, Phoenix, AZ, United states, 2094\u20132100","DOI":"10.1609\/aaai.v30i1.10295"},{"key":"2393_CR43","unstructured":"Schaul T, Quan J, Antonoglou I, Silver D (2016). Prioritized experience replay. In: 4th International Conference on Learning Representations, ICLR 2016, San Juan, Puerto Rico"},{"issue":"11","key":"2393_CR44","doi-asserted-by":"publisher","first-page":"1267","DOI":"10.3390\/jmse9111267","volume":"9","author":"Z Zhu","year":"2021","unstructured":"Zhu Z, Hu C, Zhu C, Zhu Y, Sheng Y (2021) An improved dueling deep double-q network based on prioritized experience replay for path planning of unmanned surface vehicles. Journal of Marine Science and Engineering 9(11):1267","journal-title":"Journal of Marine Science and Engineering"}],"container-title":["International Journal of Machine Learning and Cybernetics"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s13042-024-02393-z.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s13042-024-02393-z\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s13042-024-02393-z.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,4,7]],"date-time":"2025-04-07T07:26:24Z","timestamp":1744010784000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s13042-024-02393-z"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2024,10,10]]},"references-count":44,"journal-issue":{"issue":"4","published-print":{"date-parts":[[2025,4]]}},"alternative-id":["2393"],"URL":"https:\/\/doi.org\/10.1007\/s13042-024-02393-z","relation":{},"ISSN":["1868-8071","1868-808X"],"issn-type":[{"value":"1868-8071","type":"print"},{"value":"1868-808X","type":"electronic"}],"subject":[],"published":{"date-parts":[[2024,10,10]]},"assertion":[{"value":"22 April 2024","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"20 September 2024","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"10 October 2024","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}]}}