{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,28]],"date-time":"2026-07-28T10:51:11Z","timestamp":1785235871546,"version":"3.55.0"},"reference-count":46,"publisher":"Elsevier BV","license":[{"start":{"date-parts":[[2026,7,1]],"date-time":"2026-07-01T00:00:00Z","timestamp":1782864000000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.elsevier.com\/tdm\/userlicense\/1.0\/"},{"start":{"date-parts":[[2026,7,1]],"date-time":"2026-07-01T00:00:00Z","timestamp":1782864000000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.elsevier.com\/legal\/tdmrep-license"},{"start":{"date-parts":[[2026,7,1]],"date-time":"2026-07-01T00:00:00Z","timestamp":1782864000000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-017"},{"start":{"date-parts":[[2026,7,1]],"date-time":"2026-07-01T00:00:00Z","timestamp":1782864000000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-037"},{"start":{"date-parts":[[2026,7,1]],"date-time":"2026-07-01T00:00:00Z","timestamp":1782864000000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-012"},{"start":{"date-parts":[[2026,7,1]],"date-time":"2026-07-01T00:00:00Z","timestamp":1782864000000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-029"},{"start":{"date-parts":[[2026,7,1]],"date-time":"2026-07-01T00:00:00Z","timestamp":1782864000000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-004"}],"funder":[{"DOI":"10.13039\/501100012226","name":"Fundamental Research Funds for the Central Universities","doi-asserted-by":"publisher","id":[{"id":"10.13039\/501100012226","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/501100013804","name":"Fundamental Research Funds for the Central Universities","doi-asserted-by":"publisher","id":[{"id":"10.13039\/501100013804","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["elsevier.com","sciencedirect.com"],"crossmark-restriction":true},"short-container-title":["Expert Systems with Applications"],"published-print":{"date-parts":[[2026,7]]},"DOI":"10.1016\/j.eswa.2026.132289","type":"journal-article","created":{"date-parts":[[2026,3,31]],"date-time":"2026-03-31T15:58:06Z","timestamp":1774972686000},"page":"132289","update-policy":"https:\/\/doi.org\/10.1016\/elsevier_cm_policy","source":"Crossref","is-referenced-by-count":1,"special_numbering":"C","title":["Structural entropy guided hierarchical symmetric multi-agent reinforcement learning"],"prefix":"10.1016","volume":"321","author":[{"ORCID":"https:\/\/orcid.org\/0000-0003-0802-8633","authenticated-orcid":false,"given":"Yongkai","family":"Tian","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-3354-8625","authenticated-orcid":false,"given":"Xin","family":"Yu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0009-0006-1325-3009","authenticated-orcid":false,"given":"Yirong","family":"Qi","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0009-0002-1094-6927","authenticated-orcid":false,"given":"Li","family":"Wang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-6219-1741","authenticated-orcid":false,"given":"Pu","family":"Feng","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0003-2998-8828","authenticated-orcid":false,"given":"Wenjun","family":"Wu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0003-4298-9358","authenticated-orcid":false,"given":"Rongye","family":"Shi","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-4157-9931","authenticated-orcid":false,"given":"Jie","family":"Luo","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"78","reference":[{"key":"10.1016\/j.eswa.2026.132289_bib0001","series-title":"International conference on learning representations","article-title":"Geometric and physical quantities improve e (3) equivariant message passing","author":"Brandstetter","year":"2022"},{"key":"10.1016\/j.eswa.2026.132289_bib0002","doi-asserted-by":"crossref","DOI":"10.1016\/j.eswa.2025.126938","article-title":"Xlight: An interpretable multi-agent reinforcement learning approach for traffic signal control","volume":"273","author":"Cai","year":"2025","journal-title":"Expert Systems with Applications"},{"key":"10.1016\/j.eswa.2026.132289_bib0003","series-title":"International conference on machine learning","first-page":"7640","article-title":"E(3)-Equivariant actor-critic methods for cooperative multi-agent reinforcement learning","author":"Chen","year":"2024"},{"key":"10.1016\/j.eswa.2026.132289_bib0004","series-title":"2024 international joint conference on neural networks (IJCNN)","first-page":"1","article-title":"Sgcd: Subgroup contribution decomposition for multi-agent reinforcement learning","author":"Chen","year":"2024"},{"key":"10.1016\/j.eswa.2026.132289_bib0005","series-title":"Thirty-seventh conference on neural information processing systems datasets and benchmarks track","article-title":"SMACv2: An improved benchmark for cooperative multi-agent reinforcement learning","author":"Ellis","year":"2023"},{"key":"10.1016\/j.eswa.2026.132289_bib0006","article-title":"Self-clustering hierarchical multi-agent reinforcement learning with extensible cooperation graph","author":"Fu","year":"2024","journal-title":"IEEE Transactions on Emerging Topics in Computational Intelligence"},{"key":"10.1016\/j.eswa.2026.132289_bib0007","series-title":"2024\u202fIEEE\/RSJ international conference on intelligent robots and systems (IROS)","first-page":"5732","article-title":"Graph neural network-based multi-agent reinforcement learning for resilient distributed coordination of multi-robot systems","author":"Goeckner","year":"2024"},{"key":"10.1016\/j.eswa.2026.132289_bib0008","series-title":"Deep learning","volume":"vol. 1","author":"Goodfellow","year":"2016"},{"issue":"2","key":"10.1016\/j.eswa.2026.132289_bib0009","doi-asserted-by":"crossref","first-page":"895","DOI":"10.1007\/s10462-021-09996-w","article-title":"Multi-agent deep reinforcement learning: A survey","volume":"55","author":"Gronauer","year":"2022","journal-title":"Artificial Intelligence Review"},{"key":"10.1016\/j.eswa.2026.132289_bib0010","doi-asserted-by":"crossref","DOI":"10.1016\/j.neunet.2025.107554","article-title":"Learning to solve combinatorial optimization problems with heterophily","author":"Guo","year":"2025","journal-title":"Neural Networks"},{"key":"10.1016\/j.eswa.2026.132289_bib0011","series-title":"The eleventh international conference on learning representations.","article-title":"Boosting multi-agent reinforcement learning via permutation invariant and permutation equivariant networks","author":"Hao","year":"2023"},{"issue":"1","key":"10.1016\/j.eswa.2026.132289_bib0012","doi-asserted-by":"crossref","first-page":"172","DOI":"10.3390\/make4010009","article-title":"Hierarchical reinforcement learning: A survey and open research challenges","volume":"4","author":"Hutsebaut-Buysse","year":"2022","journal-title":"Machine Learning and Knowledge Extraction"},{"issue":"54","key":"10.1016\/j.eswa.2026.132289_bib0013","first-page":"1","article-title":"Deep reinforcement learning for swarm systems","volume":"20","author":"H\u00fcttenrauch","year":"2019","journal-title":"Journal of Machine Learning Research"},{"key":"10.1016\/j.eswa.2026.132289_bib0014","doi-asserted-by":"crossref","DOI":"10.1016\/j.eswa.2024.124627","article-title":"Optimizing gate control coordination signal for urban traffic network boundaries using multi-agent deep reinforcement learning","volume":"255","author":"Kang","year":"2024","journal-title":"Expert Systems with Applications"},{"issue":"6","key":"10.1016\/j.eswa.2026.132289_bib0015","doi-asserted-by":"crossref","first-page":"422","DOI":"10.1038\/s42254-021-00314-5","article-title":"Physics-informed machine learning","volume":"3","author":"Karniadakis","year":"2021","journal-title":"Nature Reviews Physics"},{"key":"10.1016\/j.eswa.2026.132289_bib0016","doi-asserted-by":"crossref","DOI":"10.1016\/j.eswa.2025.127856","article-title":"Heterogeneous multi-agent reinforcement learning based on modularized policy network","author":"Kim","year":"2025","journal-title":"Expert Systems with Applications"},{"issue":"6","key":"10.1016\/j.eswa.2026.132289_bib0017","doi-asserted-by":"crossref","first-page":"3290","DOI":"10.1109\/TIT.2016.2555904","article-title":"Structural information and dynamical complexity of networks","volume":"62","author":"Li","year":"2016","journal-title":"IEEE Transactions on Information Theory"},{"issue":"1","key":"10.1016\/j.eswa.2026.132289_bib0018","doi-asserted-by":"crossref","first-page":"3265","DOI":"10.1038\/s41467-018-05691-7","article-title":"Decoding topologically associating domains with ultra-low resolution hi-c data by graph structural entropy","volume":"9","author":"Li","year":"2018","journal-title":"Nature Communications"},{"key":"10.1016\/j.eswa.2026.132289_bib0019","doi-asserted-by":"crossref","first-page":"211","DOI":"10.1016\/j.physa.2016.09.009","article-title":"Resistance maximization principle for defending networks against virus attack","volume":"466","author":"Li","year":"2017","journal-title":"Physica A: Statistical Mechanics and its Applications"},{"key":"10.1016\/j.eswa.2026.132289_bib0020","unstructured":"Li, Y., Wang, L., Yang, J., Wang, E., Wang, Z., Zhao, T., & Zha, H. (2021). Permutation invariant policy optimization for mean-field multi-agent reinforcement learning: A principled approach. arXiv preprint arXiv: 2105.08268."},{"key":"10.1016\/j.eswa.2026.132289_bib0021","series-title":"Machine learning proceedings 1994","first-page":"157","article-title":"Markov games as a framework for multi-agent reinforcement learning","author":"Littman","year":"1994"},{"key":"10.1016\/j.eswa.2026.132289_bib0022","series-title":"Conference on robot learning","first-page":"590","article-title":"Pic: permutation invariant critic for multi-agent deep reinforcement learning","author":"Liu","year":"2020"},{"key":"10.1016\/j.eswa.2026.132289_bib0023","article-title":"Vehicle-level fairness-oriented constrained multi-agent reinforcement learning for adaptive traffic signal control","author":"Liu","year":"2025","journal-title":"IEEE Transactions on Intelligent Transportation Systems"},{"key":"10.1016\/j.eswa.2026.132289_bib0024","series-title":"The thirty-eighth annual conference on neural information processing systems","article-title":"Boosting sample efficiency and generalization in multi-agent reinforcement learning via equivariance","author":"McClellan","year":"2024"},{"key":"10.1016\/j.eswa.2026.132289_bib0025","series-title":"International conference on machine learning","first-page":"25817","article-title":"Scalable multi-agent reinforcement learning through intelligent information aggregation","author":"Nayak","year":"2023"},{"issue":"2","key":"10.1016\/j.eswa.2026.132289_bib0026","doi-asserted-by":"crossref","first-page":"73","DOI":"10.1016\/j.jai.2024.02.003","article-title":"A survey on multi-agent reinforcement learning and its application","volume":"3","author":"Ning","year":"2024","journal-title":"Journal of Automation and Intelligence"},{"key":"10.1016\/j.eswa.2026.132289_bib0027","series-title":"International encyclopedia of human geography (second edition)","first-page":"415","article-title":"Nonlinear dynamic spatial systems","author":"O\u2019Sullivan","year":"2009"},{"key":"10.1016\/j.eswa.2026.132289_bib0028","doi-asserted-by":"crossref","DOI":"10.1109\/TCYB.2024.3453892","article-title":"Enhancing collaboration in heterogeneous multiagent systems through communication complementary graph","author":"Peng","year":"2024","journal-title":"IEEE Transactions on Cybernetics"},{"key":"10.1016\/j.eswa.2026.132289_bib0029","series-title":"International conference on learning representations","article-title":"Multi-agent MDP homomorphic networks","author":"van der Pol","year":"2022"},{"key":"10.1016\/j.eswa.2026.132289_bib0030","first-page":"4199","article-title":"Mdp homomorphic networks: Group symmetries in reinforcement learning","volume":"33","author":"Van der Pol","year":"2020","journal-title":"Advances in Neural Information Processing Systems"},{"key":"10.1016\/j.eswa.2026.132289_bib0031","doi-asserted-by":"crossref","first-page":"686","DOI":"10.1016\/j.jcp.2018.10.045","article-title":"Physics-informed neural networks: A deep learning framework for solving forward and inverse problems involving nonlinear partial differential equations","volume":"378","author":"Raissi","year":"2019","journal-title":"Journal of Computational Physics"},{"key":"10.1016\/j.eswa.2026.132289_bib0032","series-title":"Proceedings of the AAAI conference on artificial intelligence","first-page":"7236","article-title":"Multi-agent actor-critic with hierarchical graph attention network","volume":"vol. 34","author":"Ryu","year":"2020"},{"key":"10.1016\/j.eswa.2026.132289_bib0033","article-title":"The starcraft multi-agent challenge","volume":"abs\/1902.04043","author":"Samvelyan","year":"2019","journal-title":"CoRR"},{"key":"10.1016\/j.eswa.2026.132289_bib0034","series-title":"International conference on machine learning","first-page":"9323","article-title":"E (n) equivariant graph neural networks","author":"Satorras","year":"2021"},{"key":"10.1016\/j.eswa.2026.132289_bib0035","doi-asserted-by":"crossref","DOI":"10.1109\/TMC.2025.3553285","article-title":"Symmetry-informed MARL: A decentralized and cooperative UAV swarm control approach for communication coverage","author":"Shi","year":"2025","journal-title":"IEEE Transactions on Mobile Computing"},{"key":"10.1016\/j.eswa.2026.132289_bib0036","series-title":"Ecai 2024","first-page":"2202","article-title":"Exploiting hierarchical symmetry in multi-agent reinforcement learning","author":"Tian","year":"2024"},{"key":"10.1016\/j.eswa.2026.132289_bib0037","series-title":"International conference on learning representations","article-title":"SO(2)-Equivariant reinforcement learning","author":"Wang","year":"2022"},{"key":"10.1016\/j.eswa.2026.132289_bib0038","series-title":"International conference on machine learning","first-page":"24017","article-title":"Structural entropy guided graph hierarchical pooling","author":"Wu","year":"2022"},{"key":"10.1016\/j.eswa.2026.132289_bib0039","article-title":"Roco: Role-oriented communication for efficient multi-agent reinforcement learning","author":"Xie","year":"2025","journal-title":"Expert Systems with Applications"},{"key":"10.1016\/j.eswa.2026.132289_bib0040","series-title":"Proceedings of the AAAI conference on artificial intelligence","first-page":"11744","article-title":"Hierarchical mean-field deep reinforcement learning for large-scale multiagent systems","volume":"vol. 37","author":"Yu","year":"2023"},{"key":"10.1016\/j.eswa.2026.132289_bib0041","first-page":"24611","article-title":"The surprising effectiveness of ppo in cooperative multi-agent games","volume":"35","author":"Yu","year":"2022","journal-title":"Advances in Neural Information Processing Systems"},{"key":"10.1016\/j.eswa.2026.132289_bib0042","series-title":"Proceedings of the AAAI conference on artificial intelligence","first-page":"17583","article-title":"Leveraging partial symmetry for multi-agent reinforcement learning","volume":"vol. 38","author":"Yu","year":"2024"},{"key":"10.1016\/j.eswa.2026.132289_bib0043","series-title":"Ecai 2023","first-page":"2946","article-title":"Esp: Exploiting symmetry prior for multi-agent reinforcement learning","author":"Yu","year":"2023"},{"key":"10.1016\/j.eswa.2026.132289_bib0044","series-title":"2024\u202fIEEE international conference on robotics and automation (ICRA)","first-page":"10814","article-title":"Adaptaug: Adaptive data augmentation framework for multi-agent reinforcement learning","author":"Yu","year":"2024"},{"key":"10.1016\/j.eswa.2026.132289_bib0045","doi-asserted-by":"crossref","first-page":"1","DOI":"10.1186\/s13059-020-02234-6","article-title":"SuperTAD: Robust detection of hierarchical topologically associated domains with optimized structural information","volume":"22","author":"Zhang","year":"2021","journal-title":"Genome Biology"},{"key":"10.1016\/j.eswa.2026.132289_bib0046","article-title":"Efficient deployment of multiple jumping robots in uneven terrains using deep reinforcement learning","author":"Zhou","year":"2025","journal-title":"Expert Systems with Applications"}],"container-title":["Expert Systems with Applications"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/api.elsevier.com\/content\/article\/PII:S0957417426012029?httpAccept=text\/xml","content-type":"text\/xml","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/api.elsevier.com\/content\/article\/PII:S0957417426012029?httpAccept=text\/plain","content-type":"text\/plain","content-version":"vor","intended-application":"text-mining"}],"deposited":{"date-parts":[[2026,6,8]],"date-time":"2026-06-08T16:03:53Z","timestamp":1780934633000},"score":1,"resource":{"primary":{"URL":"https:\/\/linkinghub.elsevier.com\/retrieve\/pii\/S0957417426012029"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,7]]},"references-count":46,"alternative-id":["S0957417426012029"],"URL":"https:\/\/doi.org\/10.1016\/j.eswa.2026.132289","relation":{},"ISSN":["0957-4174"],"issn-type":[{"value":"0957-4174","type":"print"}],"subject":[],"published":{"date-parts":[[2026,7]]},"assertion":[{"value":"Elsevier","name":"publisher","label":"This article is maintained by"},{"value":"Structural entropy guided hierarchical symmetric multi-agent reinforcement learning","name":"articletitle","label":"Article Title"},{"value":"Expert Systems with Applications","name":"journaltitle","label":"Journal Title"},{"value":"https:\/\/doi.org\/10.1016\/j.eswa.2026.132289","name":"articlelink","label":"CrossRef DOI link to publisher maintained version"},{"value":"article","name":"content_type","label":"Content Type"},{"value":"\u00a9 2026 Elsevier Ltd. All rights are reserved, including those for text and data mining, AI training, and similar technologies.","name":"copyright","label":"Copyright"}],"article-number":"132289"}}