{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,4,23]],"date-time":"2026-04-23T00:47:56Z","timestamp":1776905276155,"version":"3.51.2"},"reference-count":95,"publisher":"Springer Science and Business Media LLC","issue":"6","license":[{"start":{"date-parts":[[2025,4,23]],"date-time":"2025-04-23T00:00:00Z","timestamp":1745366400000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2025,4,23]],"date-time":"2025-04-23T00:00:00Z","timestamp":1745366400000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Mach Learn"],"published-print":{"date-parts":[[2025,6]]},"DOI":"10.1007\/s10994-025-06764-7","type":"journal-article","created":{"date-parts":[[2025,4,23]],"date-time":"2025-04-23T17:41:00Z","timestamp":1745430060000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["Inductive biases for zero-shot systematic generalization in language-informed reinforcement learning"],"prefix":"10.1007","volume":"114","author":[{"given":"Negin Hashemi","family":"Dijujin","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Seyed Roozbeh Razavi","family":"Rohani","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Mohammad Mahdi","family":"Samiei","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Mahdieh Soleymani","family":"Baghshah","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2025,4,23]]},"reference":[{"key":"6764_CR1","unstructured":"Abdin, M., Aneja, J., Awadalla, H., Awadallah, A., Awan, A.A., & Bach, N. (2024). Phi-3 technical report: A highly capable language model locally on your phone. arXiv preprint arXiv:2404.14219,"},{"key":"6764_CR2","unstructured":"Agarwal, R., Machado, M.C., Castro, P.S., & Bellemare, M.G. (2021). Contrastive behavioral similarity embeddings for generalization in reinforcement learning. In International conference on learning representations."},{"key":"6764_CR3","unstructured":"Ahn, M., Brohan, A., Brown, N., Chebotar, Y., Cortes, O., & David, B. (2022). Do as I can, not as i say: Grounding language in robotic affordances. arXiv preprint arXiv:2204.01691,"},{"key":"6764_CR4","unstructured":"Akakzia, A., Colas, C., Oudeyer, P- Y., Chetouani, M., & Sigaud, O. (2021). Grounding language to autonomously-acquired skills via goal generation. In International conference on learning representations."},{"key":"6764_CR5","unstructured":"Akkaya, I., Andrychowicz, M., Chociej, M., Litwin, M., McGrew, B., & Petron, A. (2019). Solving Rubik\u2019s cube with a robot hand. arXiv preprint arXiv:1910.07113,"},{"key":"6764_CR6","unstructured":"Aky\u00fcrek, E., Aky\u00fcrek, A.F., & Andreas, J. (2021). Learning to recombine and resample data for compositional generalization. In International conference on learning representations."},{"key":"6764_CR7","unstructured":"Alias Parth Goyal, A.G., Didolkar, A., Ke, N.R., Blundell, C., Beaudoin, P., Heess, N., & Bengio, Y. (2021). Neural production systems. Advances in Neural Information Processing Systems, 34, 25673\u201325687."},{"key":"6764_CR8","unstructured":"Bahdanau, D., Hill, F., Leike, J., Hughes, E., Kohli, P., & Grefenstette, E. (2019). Learning to understand goal specifications by modelling reward. In International conference on learning representations."},{"key":"6764_CR9","unstructured":"Bahdanau, D., Murty, S., Noukhovitch, M., Nguyen, T.H., de Vries, H., & Courville, A. (2019). Systematic generalization: What is required and can it be learned? In International conference on learning representations."},{"key":"6764_CR10","doi-asserted-by":"crossref","unstructured":"Booch, G., Fabiano, F., Horesh, L., Kate, K., Lenchner, J., & Linck, N. (2021). Thinking fast and slow in AI. In Proceedings of the AAAI conference on artificial intelligence (Vol.\u00a035, pp. 15042\u201315046).","DOI":"10.1609\/aaai.v35i17.17765"},{"key":"6764_CR11","doi-asserted-by":"crossref","unstructured":"Brohan, A., Brown, N., Carbajal, J., Chebotar, Y., Dabis, J., & Finn, C. (2022). Rt-1: Robotics transformer for real-world control at scale. arXiv preprint arXiv:2212.06817,","DOI":"10.15607\/RSS.2023.XIX.025"},{"key":"6764_CR12","doi-asserted-by":"crossref","unstructured":"Calabretta, R., & Parisi, D. (2005). Evolutionary connectionism and mind\/brain modularity. Modularity, 309.","DOI":"10.7551\/mitpress\/4734.003.0025"},{"key":"6764_CR13","unstructured":"Campero, A., Raileanu, R., Kuttler, H., Tenenbaum, J.B., Rockt\u00e4schel, T., & Grefenstette, E. (2021). Learning with amigo: Adversarially motivated intrinsic goals. In International conference on learning representations."},{"key":"6764_CR14","unstructured":"Cao, T., Wang, J., Zhang, Y., & Manivasagam, S. (2020). Babyai++: Towards grounded-language learning beyond memorization. arXiv preprint arXiv:2004.07200, ,"},{"key":"6764_CR15","doi-asserted-by":"crossref","unstructured":"Cao, Y., Zhao, H., Cheng, Y., Shu, T., Chen, Y., Liu, G., & Li, Y. (2024). Survey on large language model-enhanced reinforcement learning: Concept, taxonomy, and methods. IEEE Transactions on Neural Networks and Learning Systems.","DOI":"10.1109\/TNNLS.2024.3497992"},{"key":"6764_CR16","doi-asserted-by":"crossref","unstructured":"Carta, T., Oudeyer, P- Y., Sigaud, O., & Lamprier, S. (2022). EAGER: Asking and answering questions for automatic reward shaping in language-guided. In A.H.\u00a0Oh, A.\u00a0Agarwal, D.\u00a0Belgrave, and K.\u00a0Cho (Eds.), Advances in neural information processing systems.","DOI":"10.52202\/068431-0906"},{"key":"6764_CR17","unstructured":"Carta, T., Romac, C., Wolf, T., Lamprier, S., Sigaud, O., & Oudeyer, P- Y. (2023). Grounding large language models in interactive environments with online reinforcement learning. In International conference on machine learning (pp. 3676\u20133713)."},{"key":"6764_CR18","unstructured":"Chen, V., Gupta, A., & Marino, K. (2020). Ask your humans: Using human instructions to improve generalization in reinforcement learning. arXiv preprint arXiv:2011.00517,"},{"key":"6764_CR19","unstructured":"Chen, X., Liang, C., Yu, A.W., Song, D., & Zhou, D. (2020). Compositional Generalization via Neural-Symbolic Stack Machines. In Advances in neural information processing systems (Vol.\u00a033, pp. 1690\u20131701). Curran Associates, Inc."},{"key":"6764_CR20","unstructured":"Chevalier-Boisvert, M., Bahdanau, D., Lahlou, S., Willems, L., Saharia, C., Nguyen, T.H., & Bengio, Y. (2019). BabyAI: First steps towards grounded language learning with a human in the loop. In International conference on learning representations."},{"key":"6764_CR21","unstructured":"Cobbe, K., Klimov, O., Hesse, C., Kim, T., & Schulman, J. (2019). Quantifying generalization in reinforcement learning. In K.\u00a0Chaudhuri and R.\u00a0Salakhutdinov (Eds.), Proceedings of the 36th international conference on machine learning (Vol.\u00a097, pp. 1282\u20131289). PMLR."},{"key":"6764_CR22","unstructured":"Colas, C., Karch, T., Lair, N., Dussoux, J- M., Moulin-Frier, C., Dominey, P., & Oudeyer, P- Y. (2020). Language as a cognitive tool to imagine goals in curiosity driven exploration. In Advances in neural information processing systems (Vol.\u00a033, pp. 3761\u20133774). Curran Associates, Inc."},{"key":"6764_CR23","unstructured":"Co-Reyes, J.D., Gupta, A., Sanjeev, S., Altieri, N., DeNero, J., Abbeel, P., & Levine, S. (2019). Meta-learning language-guided policy learning. In International conference on learning representations."},{"key":"6764_CR24","doi-asserted-by":"crossref","unstructured":"C\u00f4t\u00e9, M- A., K\u00e1d\u00e1r, A., Yuan, X., Kybartas, B., Barnes, T., & Fine, E. (2019). TextWorld: A learning environment for text-based games. In Computer games: 7th workshop, CGW 2018, held in conjunction with the 27th international conference on artificial intelligence, IJCAI 2018, Stockholm, Sweden, July 13, 2018, revised selected papers 7 (pp. 41\u201375).","DOI":"10.1007\/978-3-030-24337-1_3"},{"key":"6764_CR25","doi-asserted-by":"crossref","unstructured":"Driscoll, L., Shenoy, K., & Sussillo, D. (2022). Flexible multitask computation in recurrent networks utilizes shared dynamical motifs. bioRxiv, 2022\u201308,","DOI":"10.1101\/2022.08.15.503870"},{"key":"6764_CR26","unstructured":"Duan, Y., Schulman, J., Chen, X., Bartlett, P.L., Sutskever, I., & Abbeel, P. (2016). Rl2: Fast reinforcement learning via slow reinforcement learning. arXiv preprint arXiv:1611.02779,"},{"issue":"9","key":"6764_CR27","doi-asserted-by":"publisher","first-page":"547","DOI":"10.1038\/nrn.2017.74","volume":"18","author":"H Eichenbaum","year":"2017","unstructured":"Eichenbaum, H. (2017). Prefrontal-hippocampal interactions in episodic memory. Nature Reviews Neuroscience, 18(9), 547\u2013558.","journal-title":"Nature Reviews Neuroscience"},{"issue":"4","key":"6764_CR28","doi-asserted-by":"publisher","first-page":"e1007720","DOI":"10.1371\/journal.pcbi.1007720","volume":"16","author":"NT Franklin","year":"2020","unstructured":"Franklin, N. T., & Frank, M. J. (2020). Generalizing to generalize: Humans flexibly switch between compositional and conjunctive structures during reinforcement learning. PLoS Computational Biology, 16(4), e1007720.","journal-title":"PLoS Computational Biology"},{"key":"6764_CR29","unstructured":"Fu, J., Korattikara, A., Levine, S., & Guadarrama, S. (2019). From language to goals: Inverse reinforcement learning for vision-based instruction following. In International conference on learning representations."},{"key":"6764_CR30","doi-asserted-by":"crossref","unstructured":"Geffner, H. (2022). Target Languages (vs. Inductive Biases) for Learning to Act and Plan. In Proceedings of the AAAI conference on artificial intelligence (Vol.\u00a036, pp. 12326\u201312333). (Issue: 11)","DOI":"10.1609\/aaai.v36i11.21497"},{"key":"6764_CR31","doi-asserted-by":"publisher","first-page":"104059","DOI":"10.1016\/j.cognition.2019.104059","volume":"194","author":"C Gonz\u00e1lez-Garc\u00eda","year":"2020","unstructured":"Gonz\u00e1lez-Garc\u00eda, C., Formica, S., Liefooghe, B., & Brass, M. (2020). Attentional prioritization reconfigures novel instructions into action-oriented task sets. Cognition, 194, 104059.","journal-title":"Cognition"},{"key":"6764_CR32","doi-asserted-by":"publisher","first-page":"117608","DOI":"10.1016\/j.neuroimage.2020.117608","volume":"226","author":"C Gonz\u00e1lez-Garc\u00eda","year":"2021","unstructured":"Gonz\u00e1lez-Garc\u00eda, C., Formica, S., Wisniewski, D., & Brass, M. (2021). Frontoparietal action-oriented codes support novel instruction implementation. NeuroImage, 226, 117608.","journal-title":"NeuroImage"},{"key":"6764_CR33","doi-asserted-by":"crossref","unstructured":"Goyal, P., Niekum, S., & Mooney, R.J. (2019). Using natural language for reward shaping in reinforcement learning. In Proceedings of the twenty-eighth international joint conference on artificial intelligence, IJCAI-19 (pp. 2385\u20132391). International Joint Conferences on Artificial Intelligence Organization.","DOI":"10.24963\/ijcai.2019\/331"},{"key":"6764_CR34","doi-asserted-by":"crossref","unstructured":"Hein, A., & Diepold, K. (2022). A minimal model for compositional generalization on gSCAN. In Proceedings of the Fifth BlackboxNLP Workshop on Analyzing and Interpreting Neural Networks for NLP (pp. 1\u201315). Abu Dhabi, United Arab Emirates (Hybrid): Association for Computational Linguistics.","DOI":"10.18653\/v1\/2022.blackboxnlp-1.1"},{"key":"6764_CR35","doi-asserted-by":"crossref","unstructured":"Hejna, J., Abbeel, P., & Pinto, L. (2023). Improving long-horizon imitation through instruction prediction. In Proceedings of the aaai conference on artificial intelligence (Vol.\u00a037, pp. 7857\u20137865).","DOI":"10.1609\/aaai.v37i7.25951"},{"key":"6764_CR36","unstructured":"Hill, F., Mokra, S., Wong, N., & Harley, T. (2020). Human instruction-following with deep reinforcement learning via transfer-learning from text. arXiv preprint arXiv:2005.09382"},{"key":"6764_CR37","unstructured":"Huang, W., Xia, F., Xiao, T., Chan, H., Liang, J., & Florence, P. (2022). Inner monologue: Embodied reasoning through planning with language models. arXiv preprint arXiv:2207.05608,"},{"key":"6764_CR38","first-page":"32225","volume":"35","author":"T Ito","year":"2022","unstructured":"Ito, T., Klinger, T., Schultz, D., Murray, J., Cole, M., & Rigotti, M. (2022). Compositional generalization through abstract representations in human and artificial neural networks. Advances in Neural Information Processing Systems, 35, 32225\u201332239.","journal-title":"Advances in Neural Information Processing Systems"},{"key":"6764_CR39","first-page":"9419","volume":"32","author":"Y Jiang","year":"2019","unstructured":"Jiang, Y., Gu, S. S., Murphy, K. P., & Finn, C. (2019). Language as an abstraction for hierarchical deep reinforcement learning. Advances in Neural Information Processing Systems, 32, 9419\u20139431.","journal-title":"Advances in Neural Information Processing Systems"},{"key":"6764_CR40","doi-asserted-by":"publisher","first-page":"119633","DOI":"10.1016\/j.neuroimage.2022.119633","volume":"263","author":"IP J\u00e4\u00e4skel\u00e4inen","year":"2022","unstructured":"J\u00e4\u00e4skel\u00e4inen, I. P., Glerean, E., Klucharev, V., Shestakova, A., & Ahveninen, J. (2022). Do sparse brain activity patterns underlie human cognition? NeuroImage, 263, 119633.","journal-title":"NeuroImage"},{"key":"6764_CR41","unstructured":"Keysers, D., Sch\u00e4rli, N., Scales, N., Buisman, H., Furrer, D., Kashubin, S., & Bousquet, O. (2020). Measuring compositional generalization: A comprehensive method on realistic data. In International conference on learning representations."},{"key":"6764_CR42","doi-asserted-by":"publisher","first-page":"201","DOI":"10.1613\/jair.1.14174","volume":"76","author":"R Kirk","year":"2023","unstructured":"Kirk, R., Zhang, A., Grefenstette, E., & Rockt\u00e4schel, T. (2023). A survey of zero-shot generalisation in deep reinforcement learning. Journal of Artificial Intelligence Research, 76, 201\u2013264.","journal-title":"Journal of Artificial Intelligence Research"},{"key":"6764_CR43","first-page":"7671","volume":"33","author":"H K\u00fcttler","year":"2020","unstructured":"K\u00fcttler, H., Nardelli, N., Miller, A., Raileanu, R., Selvatici, M., Grefenstette, E., & Rockt\u00e4schel, T. (2020). The NetHack learning environment. Advances in Neural Information Processing Systems, 33, 7671\u20137684.","journal-title":"Advances in Neural Information Processing Systems"},{"key":"6764_CR44","unstructured":"Lake, B., & Baroni, M. (2018). Generalization without systematicity: On the compositional skills of sequence-to-sequence recurrent networks. In International conference on machine learning (pp. 2873\u20132882). PMLR."},{"key":"6764_CR45","unstructured":"Lake, B.M. (2019). Compositional generalization through meta sequence-to-sequence learning. In Advances in neural information processing systems (Vol.\u00a032). Curran Associates, Inc."},{"key":"6764_CR46","unstructured":"Lake, B.M., Linzen, T., & Baroni, M. (2019). Human few-shot learning of compositional instructions. arXiv preprint arXiv:1901.04587,"},{"key":"6764_CR47","doi-asserted-by":"crossref","unstructured":"Liu, M., Zhu, M., & Zhang, W. (2022). Goal-conditioned reinforcement learning: Problems and solutions. arXiv preprint arXiv:2201.08299,","DOI":"10.24963\/ijcai.2022\/770"},{"key":"6764_CR48","first-page":"11525","volume":"33","author":"F Locatello","year":"2020","unstructured":"Locatello, F., Weissenborn, D., Unterthiner, T., Mahendran, A., Heigold, G., Uszkoreit, J., & Kipf, T. (2020). Object-centric learning with slot attention. Advances in Neural Information Processing Systems, 33, 11525\u201311538.","journal-title":"Advances in Neural Information Processing Systems"},{"key":"6764_CR49","first-page":"401","volume-title":"Thinking as a production system. The Cambridge handbook of thinking and reasoning","author":"MC Lovett","year":"2005","unstructured":"Lovett, M. C., & Anderson, J. R. (2005). Thinking as a production system. The Cambridge handbook of thinking and reasoning (pp. 401\u2013429). New York: Cambridge University Press."},{"key":"6764_CR50","unstructured":"Loynd, R., Fernandez, R., Celikyilmaz, A., Swaminathan, A., & Hausknecht, M. (2020). Working memory graphs. In International conference on machine learning (pp. 6404\u20136414)."},{"key":"6764_CR51","doi-asserted-by":"crossref","unstructured":"Luketina, J., Nardelli, N., Farquhar, G., Foerster, J., Andreas, J., Grefenstette, E., & Rockt\u00e4schel, T. (2019). A survey of reinforcement learning informed by natural language. In International joint conference on artificial intelligence,","DOI":"10.24963\/ijcai.2019\/880"},{"key":"6764_CR52","unstructured":"Madan, K., Ke, N.R., Goyal, A., Sch\u00f6lkopf, B., & Bengio, Y. (2021). Fast and slow learning of recurrent independent mechanisms. In International conference on learning representations."},{"key":"6764_CR53","first-page":"8032","volume":"34","author":"D Malik","year":"2021","unstructured":"Malik, D., Li, Y., & Ravikumar, P. (2021). When is generalizable reinforcement learning tractable? Advances in Neural Information Processing Systems, 34, 8032\u20138045.","journal-title":"Advances in Neural Information Processing Systems"},{"key":"6764_CR54","unstructured":"M\u00e1rton, C.D., Gagnon, L., Lajoie, G., & Rajan, K. (2021). Efficient and robust multi-task learning in the brain with modular latent primitives. arXiv preprint arXiv:2105.14108,"},{"key":"6764_CR55","unstructured":"Mentzer, F., Minnen, D., Agustsson, E., & Tschannen, M. (2023). Finite scalar quantization: VQ-VAE made simple. arXiv preprint  arXiv:2309.15505,"},{"key":"6764_CR56","doi-asserted-by":"publisher","first-page":"29","DOI":"10.1016\/j.cobeha.2020.07.003","volume":"38","author":"KJ Miller","year":"2021","unstructured":"Miller, K. J., & Venditto, S. J. C. (2021). Multi-step planning in the brain. Current Opinion in Behavioral Sciences, 38, 29\u201339.","journal-title":"Current Opinion in Behavioral Sciences"},{"key":"6764_CR57","first-page":"29529","volume":"34","author":"S Mirchandani","year":"2021","unstructured":"Mirchandani, S., Karamcheti, S., & Sadigh, D. (2021). Ella: Exploration through learned language abstraction. Advances in Neural Information Processing Systems, 34, 29529\u201329540.","journal-title":"Advances in Neural Information Processing Systems"},{"key":"6764_CR58","unstructured":"Mishra, N., Rohaninejad, M., Chen, X., & Abbeel, P. (2017). A simple neural attentive meta-learner. arXiv preprint arXiv:1707.03141, ,"},{"key":"6764_CR59","doi-asserted-by":"crossref","unstructured":"Misra, D., Langford, J., & Artzi, Y. (2017). Mapping instructions and visual observations to actions with reinforcement learning. In Proceedings of the 2017 conference on empirical methods in natural language processing (pp. 1004\u20131015). Association for Computational Linguistics.","DOI":"10.18653\/v1\/D17-1106"},{"issue":"3","key":"6764_CR60","first-page":"1891","volume":"27","author":"PS Muhle-Karbe","year":"2017","unstructured":"Muhle-Karbe, P. S., Duncan, J., De Baene, W., Mitchell, D. J., & Brass, M. (2017). Neural coding for instruction-based task sets in human frontoparietal and visual cortex. Cerebral Cortex, 27(3), 1891\u20131905.","journal-title":"Cerebral Cortex"},{"issue":"20","key":"6764_CR61","doi-asserted-by":"publisher","first-page":"4461","DOI":"10.1523\/JNEUROSCI.3104-20.2021","volume":"41","author":"PS Muhle-Karbe","year":"2021","unstructured":"Muhle-Karbe, P. S., Myers, N. E., & Stokes, M. G. (2021). A hierarchy of functional states in working memory. Journal of Neuroscience, 41(20), 4461\u20134475.","journal-title":"Journal of Neuroscience"},{"key":"6764_CR62","unstructured":"Nagabandi, A., Clavera, I., Liu, S., Fearing, R.S., Abbeel, P., Levine, S., & Finn, C. (2018). Learning to adapt in dynamic, real-world environments through meta-reinforcement learning. arXiv preprint arXiv:1803.11347,"},{"issue":"3","key":"6764_CR63","doi-asserted-by":"publisher","first-page":"133","DOI":"10.1038\/s42256-019-0025-4","volume":"1","author":"EO Neftci","year":"2019","unstructured":"Neftci, E. O., & Averbeck, B. B. (2019). Reinforcement learning in artificial and biological systems. Nature Machine Intelligence, 1(3), 133\u2013143.","journal-title":"Nature Machine Intelligence"},{"issue":"1","key":"6764_CR64","doi-asserted-by":"publisher","first-page":"132","DOI":"10.1016\/j.neuron.2019.08.030","volume":"104","author":"AC Nobre","year":"2019","unstructured":"Nobre, A. C., & Stokes, M. G. (2019). Premembering experience: A hierarchy of time-scales for proactive attention. Neuron, 104(1), 132\u2013146.","journal-title":"Neuron"},{"key":"6764_CR65","unstructured":"Paischer, F., Adler, T., Hofmarcher, M., & Hochreiter, S. (2023). Semantic helm: A human-readable memory for reinforcement learning. In Thirty-seventh conference on neural information processing systems."},{"key":"6764_CR66","doi-asserted-by":"publisher","first-page":"545","DOI":"10.3389\/fnins.2017.00545","volume":"11","author":"S Paneri","year":"2017","unstructured":"Paneri, S., & Gregoriou, G. G. (2017). Top-down control of visual attention by the prefrontal cortex. Functional specialization and long-range interactions. Frontiers in Neuroscience, 11, 545.","journal-title":"Frontiers in Neuroscience"},{"key":"6764_CR67","doi-asserted-by":"crossref","unstructured":"Peng, X.B., Andrychowicz, M., Zaremba, W., & Abbeel, P. (2018). Sim-to-real transfer of robotic control with dynamics randomization. In 2018 IEEE international conference on robotics and automation (ICRA) (pp. 3803\u20133810).","DOI":"10.1109\/ICRA.2018.8460528"},{"key":"6764_CR68","doi-asserted-by":"crossref","unstructured":"Perez, E., Strub, F., De\u00a0Vries, H., Dumoulin, V., & Courville, A. (2018). Film: Visual reasoning with a general conditioning layer. In Proceedings of the AAAI conference on artificial intelligence (Vol.\u00a032).","DOI":"10.1609\/aaai.v32i1.11671"},{"key":"6764_CR69","doi-asserted-by":"publisher","first-page":"146","DOI":"10.1016\/j.conb.2020.11.003","volume":"65","author":"MG Perich","year":"2020","unstructured":"Perich, M. G., & Rajan, K. (2020). Rethinking brain-wide interactions through multi-region \u2018network of networks\u2019 models. Current Opinion in Neurobiology, 65, 146\u2013151.","journal-title":"Current Opinion in Neurobiology"},{"key":"6764_CR70","unstructured":"Radford, A., Kim, J.W., Hallacy, C., Ramesh, A., Goh, G., & Agarwal, S. (2021). Learning transferable visual models from natural language supervision. In International conference on machine learning (pp. 8748\u20138763)."},{"issue":"4","key":"6764_CR71","doi-asserted-by":"publisher","first-page":"278","DOI":"10.1016\/j.tics.2019.01.010","volume":"23","author":"A Radulescu","year":"2019","unstructured":"Radulescu, A., Niv, Y., & Ballard, I. (2019). Holistic reinforcement learning: The role of structure and attention. Trends in Cognitive Sciences, 23(4), 278\u2013292.","journal-title":"Trends in Cognitive Sciences"},{"key":"6764_CR72","unstructured":"Rae, J.W., Borgeaud, S., Cai, T., Millican, K., Hoffmann, J., & Song, F. (2021). Scaling language models: Methods, analysis & insights from training gopher. arXiv preprint arXiv:2112.11446,"},{"key":"6764_CR73","doi-asserted-by":"crossref","unstructured":"Riveland, R., & Pouget, A. (2022). Generalization in sensorimotor networks configured with natural language instructions. bioRxiv, 2022\u201302,","DOI":"10.1101\/2022.02.22.481293"},{"key":"6764_CR74","doi-asserted-by":"publisher","first-page":"716671","DOI":"10.3389\/fpsyg.2021.716671","volume":"12","author":"F R\u00f6der","year":"2021","unstructured":"R\u00f6der, F., \u00d6zdemir, O., Nguyen, P. D., Wermter, S., & Eppe, M. (2021). The embodied crossmodal self forms language and interaction: a computational cognitive review. Frontiers in Psychology, 12, 716671.","journal-title":"Frontiers in Psychology"},{"key":"6764_CR75","doi-asserted-by":"crossref","unstructured":"Rohani, S.R.R., Hedayatian, S., & Baghshah, M.S. (2022). BIMRL: Brain inspired meta reinforcement learning. In 2022 IEEE\/RSJ international conference on intelligent robots and systems (IROS) (pp. 9048\u20139053).","DOI":"10.1109\/IROS47612.2022.9981250"},{"issue":"603\u2013616","key":"6764_CR76","first-page":"1","volume":"107","author":"J Russin","year":"2020","unstructured":"Russin, J., O\u2019Reilly, R. C., & Bengio, Y. (2020). Deep learning needs a prefrontal cortex. Work Bridging AI Cogn Sci, 107(603\u2013616), 1.","journal-title":"Work Bridging AI Cogn Sci"},{"key":"6764_CR77","unstructured":"Schulman, J., Wolski, F., Dhariwal, P., Radford, A., & Klimov, O. (2017). Proximal policy optimization algorithms. arXiv preprint arXiv:1707.06347,"},{"key":"6764_CR78","unstructured":"Seitzer, M., Horn, M., Zadaianchuk, A., Zietlow, D., Xiao, T., Simon-Gabriel, C- J., & Locatello, F. (2023). Bridging the gap to real-world object-centric learning. In The eleventh international conference on learning representations."},{"key":"6764_CR79","unstructured":"Shah, D., Osi\u0144ski, B., Levine, S., et\u00a0al. (2023). LM-Nav: Robotic navigation with large pre-trained models of language, vision, and action. In Conference on robot learning (pp. 492\u2013504)."},{"key":"6764_CR80","doi-asserted-by":"crossref","unstructured":"Spilsbury, S., & Ilin, A. (2022). Compositional generalization in grounded language learning via induced model sparsity. arXiv preprint arXiv:2207.02518,","DOI":"10.18653\/v1\/2022.naacl-srw.19"},{"key":"6764_CR81","doi-asserted-by":"publisher","first-page":"613","DOI":"10.1146\/annurev-psych-122414-033634","volume":"67","author":"O Sporns","year":"2016","unstructured":"Sporns, O., & Betzel, R. F. (2016). Modular brain networks. Annual Review of Psychology, 67, 613\u2013640.","journal-title":"Annual Review of Psychology"},{"key":"6764_CR82","volume-title":"Reinforcement learning: An introduction","author":"RS Sutton","year":"2018","unstructured":"Sutton, R. S., & Barto, A. G. (2018). Reinforcement learning: An introduction. MIT Press."},{"key":"6764_CR83","unstructured":"Wang, H.A., Zhong, V., Narasimhan, K., Reid, M., Zhong, V., & Zhong, V. (2021). Grounding language to entities and dynamics for generalization in reinforcement learning. In International conference on machine learning."},{"issue":"6","key":"6764_CR84","doi-asserted-by":"publisher","first-page":"860","DOI":"10.1038\/s41593-018-0147-8","volume":"21","author":"JX Wang","year":"2018","unstructured":"Wang, J. X., Kurth-Nelson, Z., Kumaran, D., Tirumala, D., Soyer, H., Leibo, J. Z., & Botvinick, M. (2018). Prefrontal cortex as a meta-reinforcement learning system. Nature Neuroscience, 21(6), 860\u2013868.","journal-title":"Nature Neuroscience"},{"key":"6764_CR85","unstructured":"Wang, R., Lehman, J., Clune, J., & Stanley, K.O. (2019). Paired open-ended trailblazer (poet): Endlessly generating increasingly complex and diverse learning environments and their solutions. arXiv preprint arXiv:1901.01753"},{"key":"6764_CR86","doi-asserted-by":"publisher","first-page":"92","DOI":"10.1016\/j.conb.2015.12.010","volume":"37","author":"XJ Wang","year":"2016","unstructured":"Wang, X. J., & Kennedy, H. (2016). Brain structure and dynamics across scales: In search of rules. Current Opinion in Neurobiology, 37, 92\u201398.","journal-title":"Current Opinion in Neurobiology"},{"key":"6764_CR87","first-page":"50932","volume":"36","author":"Z Wu","year":"2023","unstructured":"Wu, Z., Hu, J., Lu, W., Gilitschenski, I., & Garg, A. (2023). SlotDiffusion: Object-centric generative modeling with diffusion models. Advances in Neural Information Processing Systems, 36, 50932\u201350958.","journal-title":"Advances in Neural Information Processing Systems"},{"issue":"2","key":"6764_CR88","doi-asserted-by":"publisher","first-page":"297","DOI":"10.1038\/s41593-018-0310-2","volume":"22","author":"GR Yang","year":"2019","unstructured":"Yang, G. R., Joglekar, M. R., Song, H. F., Newsome, W. T., & Wang, X. J. (2019). Task representations in neural networks trained to perform many cognitive tasks. Nature Neuroscience, 22(2), 297\u2013306.","journal-title":"Nature Neuroscience"},{"key":"6764_CR89","unstructured":"Yarats, D., Kostrikov, I., & Fergus, R. (2021). Image augmentation is all you need: Regularizing deep reinforcement learning from pixels. In International conference on learning representations."},{"key":"6764_CR90","unstructured":"Zhang, A., Lyle, C., Sodhani, S., Filos, A., Kwiatkowska, M., Pineau, J., & Precup, D. (2020). Invariant causal prediction for block MDPs. In International conference on machine learning (pp. 11214\u201311224)."},{"key":"6764_CR91","unstructured":"Zhang, A., McAllister, R.T., Calandra, R., Gal, Y., & Levine, S. (2021). Learning invariant representations for reinforcement learning without reconstruction. In International conference on learning representations."},{"key":"6764_CR92","unstructured":"Zhang, H., & Guo, Y. (2022). Generalization of reinforcement learning with policy-aware adversarial data augmentation. In Decision awareness in reinforcement learning workshop at ICML 2022."},{"key":"6764_CR93","first-page":"1569","volume":"34","author":"M Zhao","year":"2021","unstructured":"Zhao, M., Liu, Z., Luan, S., Zhang, S., Precup, D., & Bengio, Y. (2021). A consciousness-inspired planning agent for model-based reinforcement learning. Advances in Neural Information Processing Systems, 34, 1569\u20131581.","journal-title":"Advances in Neural Information Processing Systems"},{"key":"6764_CR94","unstructured":"Zhong, V., Rockt\u00e4schel, T., & Grefenstette, E. (2020). RTFM: Generalising to New Environment Dynamics via Reading. In International conference on learning representations (pp. 1\u201317)."},{"issue":"1","key":"6764_CR95","first-page":"13198","volume":"22","author":"L Zintgraf","year":"2021","unstructured":"Zintgraf, L., Schulze, S., Lu, C., Feng, L., Igl, M., Shiarlis, K., & Whiteson, S. (2021). VariBAD: Variational bayes-adaptive deep RL via meta-learning. The Journal of Machine Learning Research, 22(1), 13198\u201313236.","journal-title":"The Journal of Machine Learning Research"}],"container-title":["Machine Learning"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10994-025-06764-7.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s10994-025-06764-7","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10994-025-06764-7.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,4,23]],"date-time":"2026-04-23T00:05:00Z","timestamp":1776902700000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s10994-025-06764-7"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,4,23]]},"references-count":95,"journal-issue":{"issue":"6","published-print":{"date-parts":[[2025,6]]}},"alternative-id":["6764"],"URL":"https:\/\/doi.org\/10.1007\/s10994-025-06764-7","relation":{},"ISSN":["0885-6125","1573-0565"],"issn-type":[{"value":"0885-6125","type":"print"},{"value":"1573-0565","type":"electronic"}],"subject":[],"published":{"date-parts":[[2025,4,23]]},"assertion":[{"value":"15 December 2023","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"21 January 2025","order":2,"name":"revised","label":"Revised","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"5 March 2025","order":3,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"23 April 2025","order":4,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors have no competing interests to declare that are relevant to the content of this article.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}},{"value":"Not applicable.","order":3,"name":"Ethics","group":{"name":"EthicsHeading","label":"Ethical approval"}},{"value":"Not applicable.","order":4,"name":"Ethics","group":{"name":"EthicsHeading","label":"Consent to participate"}},{"value":"Not applicable.","order":5,"name":"Ethics","group":{"name":"EthicsHeading","label":"Consent for publication"}},{"value":"This content has been made available to all.","name":"free","label":"Free to read"}],"article-number":"137"}}