{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,6,17]],"date-time":"2026-06-17T11:53:33Z","timestamp":1781697213996,"version":"3.54.5"},"reference-count":53,"publisher":"Springer Science and Business Media LLC","issue":"7","license":[{"start":{"date-parts":[[2026,6,17]],"date-time":"2026-06-17T00:00:00Z","timestamp":1781654400000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,6,17]],"date-time":"2026-06-17T00:00:00Z","timestamp":1781654400000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100008238","name":"Hebei Provincial Department of Science and Technology","doi-asserted-by":"publisher","award":["246Z0102G"],"award-info":[{"award-number":["246Z0102G"]}],"id":[{"id":"10.13039\/501100008238","id-type":"DOI","asserted-by":"publisher"}]},{"name":"Department of Science and Techonology of Zhejiang Province","award":["2025C02044"],"award-info":[{"award-number":["2025C02044"]}]},{"DOI":"10.13039\/501100003787","name":"Natural Science Foundation of Hebei Province","doi-asserted-by":"publisher","award":["F2024210008"],"award-info":[{"award-number":["F2024210008"]}],"id":[{"id":"10.13039\/501100003787","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Int. J. Mach. Learn. &amp; Cyber."],"published-print":{"date-parts":[[2026,7]]},"DOI":"10.1007\/s13042-026-03160-y","type":"journal-article","created":{"date-parts":[[2026,6,17]],"date-time":"2026-06-17T11:15:20Z","timestamp":1781694920000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["FedHetero: improving federated learning on heterogeneous data distribution through deep reinforcement learning"],"prefix":"10.1007","volume":"17","author":[{"given":"Shuo","family":"Sun","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Jiecong","family":"Wang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Jie","family":"Chen","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Xianghong","family":"Tang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Yulong","family":"Fang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Shuhai","family":"Wang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Jianwu","family":"Zheng","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Yongjian","family":"Huang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Jianfeng","family":"Wang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Wei-Tek","family":"Tsai","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Lin","family":"Liu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Xin","family":"Liu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Mingsheng","family":"Liu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2026,6,17]]},"reference":[{"key":"3160_CR1","doi-asserted-by":"publisher","first-page":"422","DOI":"10.1016\/j.future.2024.03.043","volume":"157","author":"H Liu","year":"2024","unstructured":"Liu H, Li S, Li W, Sun W (2024) Efficient decentralized optimization for edge-enabled smart manufacturing: a federated learning-based framework. Fut Gen Comput Syst 157:422\u2013435","journal-title":"Futur Gener Comput Syst"},{"key":"3160_CR2","first-page":"1087","volume":"36","author":"M Jiang","year":"2022","unstructured":"Jiang M, Wang Z, Dou Q (2022) Harmofl: harmonizing local and global drifts in federated learning on heterogeneous medical images. Proc AAAI Conf Artif Intell 36:1087\u20131095","journal-title":"Proc AAAI Conf Artif Intell"},{"key":"3160_CR3","first-page":"13509","volume":"38","author":"J Li","year":"2024","unstructured":"Li J, Liu Y, Wang W (2024) Fedns: a fast sketching newton-type algorithm for federated learning. Proc AAAI Conf Artif Intell 38:13509\u201313517","journal-title":"Proc AAAI Conf Artif Intell"},{"key":"3160_CR4","doi-asserted-by":"crossref","unstructured":"Zhang R, Mao J, Wang H, Li B, Cheng X, Yang L (2024) A survey on federated learning in intelligent transportation systems. IEEE Trans Intell Veh","DOI":"10.1109\/TIV.2024.3446319"},{"key":"3160_CR5","doi-asserted-by":"crossref","unstructured":"Vaishnav S, Khirirat S, Magn\u00fasson S (2024) Communication-adaptive gradient sparsification for federated learning with error compensation. IEEE Int Things J","DOI":"10.1109\/JIOT.2024.3490855"},{"key":"3160_CR6","doi-asserted-by":"crossref","unstructured":"Guan H, Yap P-T, Bozoki A, Liu M (2024) Federated learning for medical image analysis: a survey. Pattern Recogn, 110424","DOI":"10.1016\/j.patcog.2024.110424"},{"key":"3160_CR7","doi-asserted-by":"publisher","DOI":"10.1016\/j.artmed.2024.103024","volume":"159","author":"KN Kumar","year":"2025","unstructured":"Kumar KN, Mohan CK, Cenkeramaddi LR, Awasthi N (2025) Minimal data poisoning attack in federated learning for medical image classification: an attacker perspective. Artif Intell Med 159:103024","journal-title":"Artif Intell Med"},{"key":"3160_CR8","first-page":"1","volume":"74","author":"P Wei","year":"2025","unstructured":"Wei P, Zhou T, Liu W, Du J, Wang T, Yue G (2025) Fedpdn: personalized federated learning with inter-class similarity constraint for medical image classification through parameter decoupling. IEEE Trans Instrum Meas 74:1\u201313","journal-title":"IEEE Trans Instrum Meas"},{"key":"3160_CR9","unstructured":"Chen M, Mathews R, Ouyang T, Beaufays F (2019) Federated learning of out-of-vocabulary words. arXiv preprint arXiv:1903.10635"},{"issue":"1","key":"3160_CR10","doi-asserted-by":"publisher","first-page":"2618","DOI":"10.1109\/TCE.2023.3318754","volume":"70","author":"D Javeed","year":"2023","unstructured":"Javeed D, Saeed MS, Kumar P, Jolfaei A, Islam S, Islam AN (2023) Federated learning-based personalized recommendation systems: an overview on security and privacy challenges. IEEE Trans Consum Electron 70(1):2618\u20132627","journal-title":"IEEE Trans Consum Electron"},{"key":"3160_CR11","doi-asserted-by":"publisher","first-page":"244","DOI":"10.1016\/j.future.2022.05.003","volume":"135","author":"X Ma","year":"2022","unstructured":"Ma X, Zhu J, Lin Z, Chen S, Qin Y (2022) A state-of-the-art survey on solving non-iid data in federated learning. Futur Gener Comput Syst 135:244\u2013258","journal-title":"Futur Gener Comput Syst"},{"issue":"9","key":"3160_CR12","first-page":"3400","volume":"31","author":"F Sattler","year":"2019","unstructured":"Sattler F, Wiedemann S, M\u00fcller K-R, Samek W (2019) Robust and communication-efficient federated learning from non-iid data. IEEE TNNLS 31(9):3400\u20133413","journal-title":"IEEE TNNLS"},{"key":"3160_CR13","unstructured":"Lier S, Vries A, Herder E (2018) Robustness of federated averaging for non-iid data"},{"key":"3160_CR14","unstructured":"Zhao Y, Li M, Lai L, Suda N, Civin D, Chandra V (2018) Federated learning with non-iid data. arXiv preprint arXiv:1806.00582"},{"key":"3160_CR15","doi-asserted-by":"publisher","first-page":"3","DOI":"10.1007\/978-3-030-63076-8_1","volume":"12500","author":"L Lyul","year":"2020","unstructured":"Lyul L, Yu H, Zhao J, Yang Q (2020) Threats to federated learning. Feder Learn Priv Incent 12500:3","journal-title":"Feder Learn Priv Incent"},{"key":"3160_CR16","unstructured":"Li X, Huang K, Yang W, Wang S, Zhang Z (2019) On the convergence of fedavg on non-iid data. In: Proceedings of the ICLR"},{"issue":"10","key":"3160_CR17","first-page":"9895","volume":"7","author":"D Kwon","year":"2020","unstructured":"Kwon D, Jeon J, Park S, Kim J, Cho S (2020) Multiagent ddpg-based deep learning for smart ocean federated learning iot networks. IEEE IoT-J 7(10):9895\u20139903","journal-title":"IEEE IoT-J"},{"issue":"10","key":"3160_CR18","doi-asserted-by":"publisher","first-page":"9352","DOI":"10.1109\/TMC.2024.3365295","volume":"23","author":"Z Li","year":"2024","unstructured":"Li Z, Sun Y, Shao J, Mao Y, Wang JH, Zhang J (2024) Feature matching data synthesis for non-iid federated learning. IEEE Trans Mob Comput 23(10):9352\u20139367","journal-title":"IEEE Trans Mob Comput"},{"key":"3160_CR19","doi-asserted-by":"crossref","unstructured":"Wang H, Kaplan Z, Niu D, Li B (2020) Optimizing federated learning on non-iid data with reinforcement learning. In: Proceedings of the INFOCOM, pp 1698\u20131707. IEEE","DOI":"10.1109\/INFOCOM41043.2020.9155494"},{"issue":"7540","key":"3160_CR20","doi-asserted-by":"publisher","first-page":"529","DOI":"10.1038\/nature14236","volume":"518","author":"V Mnih","year":"2015","unstructured":"Mnih V, Kavukcuoglu K, Silver D, Rusu AA, Veness J, Bellemare MG, Graves A, Riedmiller M, Fidjeland AK, Ostrovski G et al (2015) Human-level control through deep reinforcement learning. Nature 518(7540):529\u2013533","journal-title":"Nature"},{"issue":"11","key":"3160_CR21","first-page":"2524","volume":"31","author":"L Lyu","year":"2020","unstructured":"Lyu L, Yu J, Nandakumar K, Li Y, Ma X, Jin J, Yu H, Ng KS (2020) Towards fair and privacy-preserving federated deep models. IEEE TPDS 31(11):2524\u20132541","journal-title":"IEEE TPDS"},{"key":"3160_CR22","unstructured":"Lillicrap TP, Hunt JJ, Pritzel A, Heess N, Erez T, Tassa Y, Silver D, Wierstra D (2016) Continuous control with deep reinforcement learning. In: Proceedings of the ICLR (Poster)"},{"key":"3160_CR23","unstructured":"McMahan B, Moore E, Ramage D, Hampson S, Arcas BA (2017) Communication-efficient learning of deep networks from decentralized data. In: Proceedings of the AISTATS, pp 1273\u20131282. PMLR"},{"issue":"24","key":"3160_CR24","doi-asserted-by":"publisher","first-page":"21467","DOI":"10.1109\/JIOT.2023.3298196","volume":"10","author":"G Xu","year":"2023","unstructured":"Xu G, Zhou Z, Dong J, Zhang L, Song X (2023) A blockchain-based federated learning scheme for data sharing in industrial internet of things. IEEE Int Things J 10(24):21467\u201321478. https:\/\/doi.org\/10.1109\/JIOT.2023.3298196","journal-title":"IEEE Int Things J"},{"key":"3160_CR25","doi-asserted-by":"publisher","unstructured":"Abu-Dabaseh F, Alghizzawi M, Alkhlaifat BI, Ratib\u00a0Ezmigna AA, Alzghoul A, AlSokkar AAM, Al-Gasawneh J (2024) Enhancing privacy and security in decentralized social systems: blockchain-based approach. In: 2024 2nd international conference on cyber resilience (ICCR), pp 1\u20136. https:\/\/doi.org\/10.1109\/ICCR61006.2024.10533137","DOI":"10.1109\/ICCR61006.2024.10533137"},{"issue":"4","key":"3160_CR26","doi-asserted-by":"publisher","first-page":"679","DOI":"10.3390\/electronics13040679","volume":"13","author":"C Wan","year":"2024","unstructured":"Wan C, Wang Y, Xu J, Wu J, Zhang T, Wang Y (2024) Research on privacy protection in federated learning combining distillation defense and blockchain. Electronics 13(4):679","journal-title":"Electronics"},{"key":"3160_CR27","doi-asserted-by":"crossref","unstructured":"Han B, Li B, Zhang Y, Feng P, Wolter K, Zhang H, Li Y, Jurdak R, Yuen C (2025) Repeated game-based long-term incentive mechanism for blockchain-enabled reliable federated learning in iiot. IEEE Int Things J","DOI":"10.1109\/JIOT.2025.3600225"},{"key":"3160_CR28","unstructured":"Jia R, Dao D, Wang B, Hubis FA, Hynes N, G\u00fcrel NM, Li B, Zhang C, Song D, Spanos CJ (2019) Towards efficient data valuation based on the Shapley value. In: Proceedings of the AISTATS, pp 1167\u20131176. PMLR"},{"key":"3160_CR29","doi-asserted-by":"publisher","DOI":"10.1016\/j.neucom.2025.129550","volume":"626","author":"Y Pang","year":"2025","unstructured":"Pang Y, Ni Z, Zhong X (2025) A fast federated reinforcement learning approach with phased weight-adjustment technique. Neurocomputing 626:129550","journal-title":"Neurocomputing"},{"issue":"12","key":"3160_CR30","doi-asserted-by":"publisher","first-page":"3291","DOI":"10.1109\/TPDS.2022.3150579","volume":"33","author":"Z Zhou","year":"2022","unstructured":"Zhou Z, Li Y, Ren X, Yang S (2022) Towards efficient and stable k-asynchronous federated learning with unbounded stale gradients on non-iid data. IEEE Trans Parallel Distrib Syst 33(12):3291\u20133305","journal-title":"IEEE Trans Parallel Distrib Syst"},{"issue":"8","key":"3160_CR31","doi-asserted-by":"publisher","first-page":"5168","DOI":"10.1109\/TCOMM.2021.3083316","volume":"69","author":"MK Nori","year":"2021","unstructured":"Nori MK, Yun S, Kim I-M (2021) Fast federated learning by balancing communication trade-offs. IEEE Trans Commun 69(8):5168\u20135182","journal-title":"IEEE Trans Commun"},{"key":"3160_CR32","doi-asserted-by":"publisher","DOI":"10.1016\/j.egyai.2025.100521","volume":"21","author":"J Sievers","year":"2025","unstructured":"Sievers J, Henrich P, Beichter M, Mikut R, Hagenmeyer V, Blank T, Simon F (2025) Federated reinforcement learning for sustainable and cost-efficient energy management. Energy AI 21:100521","journal-title":"Energy AI"},{"issue":"2","key":"3160_CR33","first-page":"1769","volume":"23","author":"W Sun","year":"2023","unstructured":"Sun W, Zhao Y, Ma W, Guo B, Xu L, Duong TQ (2023) Accelerating convergence of federated learning in mec with dynamic community. IEEE Trans Mob Comput 23(2):1769\u20131784","journal-title":"IEEE Trans Mob Comput"},{"key":"3160_CR34","doi-asserted-by":"crossref","unstructured":"Yang J, Zhu M, Zhou Y, Zhang Q, Ni Y (2024) A game-theoretic incentive mechanism for multi-distributor multi-agent federated learning. In: 2024 IEEE wireless communications and networking conference (WCNC), pp 1\u20135. IEEE","DOI":"10.1109\/WCNC57260.2024.10571041"},{"issue":"2","key":"3160_CR35","first-page":"389","volume":"10","author":"L Yin","year":"2024","unstructured":"Yin L, Lin S, Sun Z, Li R, He Y, Hao Z (2024) A game-theoretic approach for federated learning: a trade-off among privacy, accuracy and energy. Dig Commun Netw 10(2):389\u2013403","journal-title":"Dig Commun Netw"},{"issue":"1","key":"3160_CR36","doi-asserted-by":"publisher","first-page":"242","DOI":"10.1109\/TGCN.2024.3424552","volume":"9","author":"M Hu","year":"2024","unstructured":"Hu M, Zhang J, Wang X, Liu S, Lin Z (2024) Accelerating federated learning with model segmentation for edge networks. IEEE Trans Green Commun Netw 9(1):242\u2013254","journal-title":"IEEE Trans Green Commun Netw"},{"key":"3160_CR37","doi-asserted-by":"publisher","first-page":"14","DOI":"10.1016\/j.comcom.2020.05.037","volume":"160","author":"N Shan","year":"2020","unstructured":"Shan N, Cui X, Gao Z (2020) \u201cdrl+ fl\u2019\u2019: an intelligent resource allocation model based on deep reinforcement learning for mobile edge computing. Comput Commun 160:14\u201324","journal-title":"Comput Commun"},{"issue":"1","key":"3160_CR38","doi-asserted-by":"publisher","first-page":"904","DOI":"10.1109\/TNET.2023.3299851","volume":"32","author":"Y Liao","year":"2023","unstructured":"Liao Y, Xu Y, Xu H, Yao Z, Wang L, Qiao C (2023) Accelerating federated learning with data and model parallelism in edge computing. IEEE\/ACM Trans Netw 32(1):904\u2013918","journal-title":"IEEE\/ACM Trans Netw"},{"key":"3160_CR39","doi-asserted-by":"crossref","unstructured":"Wang K, He Q, Chen F, Jin H, Yang Y (2023) Fededge: accelerating edge-assisted federated learning. In: Proceedings of the ACM web conference 2023, pp 2895\u20132904","DOI":"10.1145\/3543507.3583264"},{"key":"3160_CR40","unstructured":"Lai F, Zhu X, Madhyastha HV, Chowdhury M (2021) Oort: efficient federated learning via guided participant selection. In: Proceedings of the OSDI, pp 19\u201335"},{"issue":"11","key":"3160_CR41","first-page":"19188","volume":"11","author":"Z Lu","year":"2024","unstructured":"Lu Z, Pan H, Dai Y, Si X, Zhang Y (2024) Federated learning with non-iid data: a survey. IEEE Int Things J 11(11):19188\u201319209","journal-title":"IEEE Int Things J"},{"key":"3160_CR42","doi-asserted-by":"crossref","unstructured":"Do V, Atif J, Lang J, Usunier N (2021) Online selection of diverse committees. In: Proceedings of the IJCAI, pp 154\u2013160","DOI":"10.24963\/ijcai.2021\/22"},{"issue":"3","key":"3160_CR43","doi-asserted-by":"publisher","first-page":"911","DOI":"10.3390\/s25030911","volume":"25","author":"AW Mamond","year":"2025","unstructured":"Mamond AW, Kundroo M, Yoo S-E, Kim S, Kim T (2025) Fldqn: cooperative multi-agent federated reinforcement learning for solving travel time minimization problems in dynamic environments using sumo simulation. Sensors 25(3):911","journal-title":"Sensors"},{"key":"3160_CR44","unstructured":"McMahan HB, Moore E, Ramage D, Arcas BA (2016) Federated learning of deep networks using model averaging. arXiv preprint arXiv:1602.05629"},{"key":"3160_CR45","doi-asserted-by":"publisher","first-page":"5693","DOI":"10.1609\/aaai.v33i01.33015693","volume":"33","author":"H Yu","year":"2019","unstructured":"Yu H, Yang S, Zhu S (2019) Parallel restarted sgd with faster convergence and less communication: demystifying why model averaging works for deep learning. Proc AAAI 33:5693\u20135700","journal-title":"Proc AAAI"},{"key":"3160_CR46","unstructured":"Su H, Chen H (2015) Experiments on parallel training of deep neural network using model averaging. arXiv preprint arXiv:1507.01239"},{"key":"3160_CR47","unstructured":"Yu C, Tang H, Renggli C, Kassing S, Singla A, Alistarh D, Zhang C, Liu J (2019) Distributed learning over unreliable networks. In: Proceedings of the ICML, pp 7202\u20137212. PMLR"},{"key":"3160_CR48","unstructured":"Watkins CJCH (1989) Learning from delayed rewards. Ph.D. thesis Kings College University of Cambridge"},{"key":"3160_CR49","unstructured":"Sutton RS, McAllester DA, Singh SP, Mansour Y et al (1999) Policy gradient methods for reinforcement learning with function approximation. In: Proceedings of the NIPS, 99, 1057\u20131063. Citeseer"},{"issue":"1","key":"3160_CR50","first-page":"1","volume":"11","author":"J Augustine","year":"2015","unstructured":"Augustine J, Chen N, Elkind E, Fanelli A, Gravin N, Shiryaev D (2015) Dynamics of profit-sharing games. Int Math 11(1):1\u201322","journal-title":"Int Math"},{"key":"3160_CR51","unstructured":"Silver D, Lever G, Heess N, Degris T, Wierstra D, Riedmiller M (2014) Deterministic policy gradient algorithms. In: Proceedings of the ICML, pp 387\u2013395. PMLR"},{"key":"3160_CR52","unstructured":"Krizhevsky A, Hinton G et al (2009) Learning multiple layers of features from tiny images. Handbook Syst Autoimm Dis 1(4)"},{"key":"3160_CR53","unstructured":"Caldas S, Duddu SMK, Wu P, Li T, Kone\u010dn\u1ef3 J, McMahan HB, Smith V, Talwalkar A (2018) Leaf: A benchmark for federated settings. arXiv preprint arXiv:1812.01097"}],"container-title":["International Journal of Machine Learning and Cybernetics"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s13042-026-03160-y.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s13042-026-03160-y","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s13042-026-03160-y.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,6,17]],"date-time":"2026-06-17T11:15:46Z","timestamp":1781694946000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s13042-026-03160-y"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,6,17]]},"references-count":53,"journal-issue":{"issue":"7","published-print":{"date-parts":[[2026,7]]}},"alternative-id":["3160"],"URL":"https:\/\/doi.org\/10.1007\/s13042-026-03160-y","relation":{},"ISSN":["1868-8071","1868-808X"],"issn-type":[{"value":"1868-8071","type":"print"},{"value":"1868-808X","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,6,17]]},"assertion":[{"value":"2 January 2026","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"4 May 2026","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"17 June 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare no conflict of interest.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}],"article-number":"352"}}