{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,30]],"date-time":"2026-07-30T14:21:04Z","timestamp":1785421264134,"version":"3.56.0"},"reference-count":244,"publisher":"Elsevier BV","issue":"5","license":[{"start":{"date-parts":[[2025,12,1]],"date-time":"2025-12-01T00:00:00Z","timestamp":1764547200000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.elsevier.com\/tdm\/userlicense\/1.0\/"},{"start":{"date-parts":[[2025,12,1]],"date-time":"2025-12-01T00:00:00Z","timestamp":1764547200000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.elsevier.com\/legal\/tdmrep-license"},{"start":{"date-parts":[[2026,4,22]],"date-time":"2026-04-22T00:00:00Z","timestamp":1776816000000},"content-version":"vor","delay-in-days":142,"URL":"http:\/\/creativecommons.org\/licenses\/by-nc-nd\/4.0\/"}],"funder":[{"DOI":"10.13039\/501100008769","name":"Julius-Maximilians-Universit\u00e4t W\u00fcrzburg","doi-asserted-by":"crossref","id":[{"id":"10.13039\/501100008769","id-type":"DOI","asserted-by":"crossref"}]}],"content-domain":{"domain":["elsevier.com","sciencedirect.com"],"crossmark-restriction":true},"short-container-title":["International Journal of Artificial Intelligence in Education"],"published-print":{"date-parts":[[2025,12]]},"DOI":"10.1007\/s40593-025-00494-6","type":"journal-article","created":{"date-parts":[[2025,7,10]],"date-time":"2025-07-10T22:11:39Z","timestamp":1752185499000},"page":"2669-2723","update-policy":"https:\/\/doi.org\/10.1016\/elsevier_cm_policy","source":"Crossref","is-referenced-by-count":26,"title":["Reinforcement Learning in Education: A Systematic Literature Review"],"prefix":"10.1016","volume":"35","author":[{"ORCID":"https:\/\/orcid.org\/0000-0002-1969-8563","authenticated-orcid":false,"given":"Anna","family":"Riedmann","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0001-5960-7921","authenticated-orcid":false,"given":"Philipp","family":"Schaper","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-2362-0080","authenticated-orcid":false,"given":"Birgit","family":"Lugrin","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"78","reference":[{"key":"10.1007\/s40593-025-00494-6_bib1","article-title":"Bridging Declarative, Procedural, and Conditional Metacognitive Knowledge Gap Using Deep Reinforcement Learning","author":"Abdelshiheed","year":"2023","journal-title":"CogSci\u201923: The 45th Annual Conference of the Cognitive Science Society."},{"key":"10.1007\/s40593-025-00494-6_bib2","series-title":"Lecture Notes in Artificial Intelligence: Vol. 13916. Artificial Intelligence in Education: 24th International Conference, AIED 2023, Tokyo, Japan, July 3\u20137, 2023, Proceedings","first-page":"291","article-title":"Leveraging deep reinforcement learning for metacognitive interventions across intelligent tutoring systems","author":"Abdelshiheed","year":"2023"},{"issue":"7","key":"10.1007\/s40593-025-00494-6_bib3","doi-asserted-by":"crossref","first-page":"1","DOI":"10.1145\/3543846","article-title":"Reinforcement learning based recommender systems: A survey","volume":"55","author":"Afsar","year":"2023","journal-title":"ACM Computing Surveys"},{"key":"10.1007\/s40593-025-00494-6_bib4","unstructured":"*Ai, F., Chen, Y., Guo, Y., Zhao, Y., Wang, Z., Fu, G., & Wang, G. Concept-aware deep knowledge tracing and exercise recommendation in an online learning system. In Proceedings of The 12th International Conference on Educational Data Mining (EDM 2019) (pp. 240\u2013245)."},{"issue":"4","key":"10.1007\/s40593-025-00494-6_bib5","doi-asserted-by":"crossref","DOI":"10.3390\/s21041292","article-title":"Reinforcement learning approaches in social robotics","volume":"21","author":"Akalin","year":"2021","journal-title":"Sensors"},{"key":"10.1007\/s40593-025-00494-6_bib6","series-title":"2021 5th International Conference on Computing Methodologies and Communication (ICCMC)","first-page":"1416","article-title":"Review on reinforcement learning, research evolution and scope of application","author":"Akanksha","year":"2021"},{"key":"10.1007\/s40593-025-00494-6_bib7","article-title":"How Much Training is Needed? Reducing Training Time using Deep Reinforcement Learning in an Intelligent Tutor","author":"Alam","year":"2024","journal-title":"Proceedings of the 17th International Conference on Educational Data Mining."},{"key":"10.1007\/s40593-025-00494-6_bib8","doi-asserted-by":"crossref","first-page":"89769","DOI":"10.1109\/ACCESS.2023.3305584","article-title":"Smart e-learning framework for personalized adaptive learning and sequential path recommendations using reinforcement learning","volume":"11","author":"Amin","year":"2023","journal-title":"IEEE Access"},{"key":"10.1007\/s40593-025-00494-6_bib9","first-page":"168","article-title":"Leveraging deep reinforcement learning for pedagogical policy induction in an intelligent tutoring system","author":"Ausin","year":"2019","journal-title":"Proceedings of The 12th International Conference on Educational Data Mining (EDM 2019)"},{"key":"10.1007\/s40593-025-00494-6_bib10","series-title":"Lecture Notes in Computer Science. ARTIFICIAL INTELLIGENCE IN EDUCATION: 22nd international conference, aied 2021","first-page":"356","article-title":"Tackling the credit assignment problem in reinforcement learning-induced pedagogical policies with neural networks","author":"Ausin","year":"2021"},{"key":"10.1007\/s40593-025-00494-6_bib11","article-title":"The impact of batch deep reinforcement learning on student performance: A simple act of explanation can go a long way","author":"Ausin","year":"2022","journal-title":"International Journal of Artificial Intelligence in Education"},{"key":"10.1007\/s40593-025-00494-6_bib12","series-title":"2020 IEEE International Conference on Robotics and Automation (ICRA)","first-page":"2746","article-title":"Deepracer: Autonomous racing platform for experimentation with sim2real reinforcement learning","author":"Balaji","year":"2020"},{"key":"10.1007\/s40593-025-00494-6_bib13","article-title":"A pilot study on logic proof tutoring using hints generated from historical student data","author":"Barnes","year":"2008","journal-title":"Proceedings of Educational Data Mining 2008, The 1st International Conference on Educational Data Mining, Montreal, Qu\u00e9bec, Canada, June 20\u201321, 2008"},{"key":"10.1007\/s40593-025-00494-6_bib14","series-title":"Lecture Notes in Computer Science: Vol. 5091. Intelligent tutoring systems: 9th international conference, ITS 2008, Montreal, Canada, June 23 - 27, 2008; proceedings","first-page":"373","article-title":"Toward automatic hint generation for logic proof tutoring using historical student data","author":"Barnes","year":"2008"},{"key":"10.1007\/s40593-025-00494-6_bib15","series-title":"ACM Digital Library, Proceedings of the 2020 CHI Conference on Human Factors in Computing Systems","first-page":"1","article-title":"Reinforcement learning for the adaptive scheduling of educational activities","author":"Bassen","year":"2020"},{"key":"10.1007\/s40593-025-00494-6_bib16","article-title":"Modeling the student with reinforcement learning","author":"Beck","year":"1997","journal-title":"Machine learning for User Modeling Workshop at the Sixth International Conference on User Modeling."},{"key":"10.1007\/s40593-025-00494-6_bib17","series-title":"Proceedings of the Fifteenth National Conference on Artificial Intelligence and Tenth Innovative Applications of Artificial Intelligence Conference, AAAI 98, IAAI 98, July 26\u201330, 1998, Madison, Wisconsin, USA","first-page":"1185","article-title":"Learning to teach with a reinforcement learning agent","author":"Beck","year":"1998"},{"key":"10.1007\/s40593-025-00494-6_bib18","first-page":"552","article-title":"ADVISOR: A machine learning architecture for intelligent tutor construction","author":"Beck","year":"2000","journal-title":"Proceedings of the Seventeenth National Conference on Artificial Intelligence and Twelfth Conference on Innovative Applications of Artificial Intelligence"},{"key":"10.1007\/s40593-025-00494-6_bib19","series-title":"Lecture Notes in Computer Science: Vol. 1839. Intelligent tutoring systems: 5th international conference; proceedings","first-page":"584","article-title":"High-level student modeling with machine learning","author":"Beck","year":"2000"},{"key":"10.1007\/s40593-025-00494-6_bib20","series-title":"Lecture Notes in Computer Science: Vol. 13355. Artificial Intelligence in Education: 23rd International Conference, AIED 2022, Durham, UK, July 27\u201331, 2022, Proceedings, Part I (1st ed. 2022, Vol. 13355","first-page":"724","article-title":"Raising student completion rates with adaptive curriculum and contextual bandits","author":"Belfer","year":"2022"},{"issue":"8","key":"10.1007\/s40593-025-00494-6_bib21","doi-asserted-by":"crossref","first-page":"716","DOI":"10.1073\/pnas.38.8.716","article-title":"On the Theory of Dynamic Programming","volume":"38","author":"Bellman","year":"1952","journal-title":"Proceedings of the National Academy of Sciences of the United States of America"},{"issue":"4","key":"10.1007\/s40593-025-00494-6_bib22","doi-asserted-by":"crossref","first-page":"264","DOI":"10.1109\/TCIAIG.2009.2035923","article-title":"Adaptive experience engine for serious games","volume":"1","author":"Bellotti","year":"2009","journal-title":"IEEE Transactions on Computational Intelligence and AI in Games"},{"issue":"1","key":"10.1007\/s40593-025-00494-6_bib23","doi-asserted-by":"crossref","first-page":"13","DOI":"10.15388\/infedu.2013.02","article-title":"Adaptive educational software by applying reinforcement learning","volume":"12","author":"Bennane","year":"2013","journal-title":"Informatics in Education"},{"key":"10.1007\/s40593-025-00494-6_bib24","first-page":"35","article-title":"Combining AI techniques into a legal agent-based intelligent tutoring system","volume":"18","author":"Bittencourt","year":"2006","journal-title":"Eighteenth International Conference on Software Engineering and Knowledge Engineering-SEKE"},{"key":"10.1007\/s40593-025-00494-6_bib25","doi-asserted-by":"crossref","first-page":"1198","DOI":"10.1016\/j.procs.2020.03.028","article-title":"Towards an adaptive e-learning system based on q-learning algorithm","volume":"170","author":"Boussakssou","year":"2020","journal-title":"Procedia Computer Science"},{"key":"10.1007\/s40593-025-00494-6_bib26","series-title":"2023 IEEE Colombian Caribbean Conference (C3)","first-page":"1","article-title":"A Formal Model for Personalized Learning Path using Artificial Intelligence for Instructional Planning with a Focus on 21st-Century Skills and Environmental Awareness","author":"Caro","year":"2023"},{"key":"10.1007\/s40593-025-00494-6_bib27","first-page":"588","article-title":"The effects of a personalized recommendation system on students\u2019 high-stakes achievement scores: A field experiment","author":"Chakraborty","year":"2021","journal-title":"Proceedings of The 14th International Conference on Educational Data Mining (EDM 2021)"},{"key":"10.1007\/s40593-025-00494-6_bib28","doi-asserted-by":"crossref","DOI":"10.1016\/j.compedu.2020.103836","article-title":"Teaching and learning with children: Impact of reciprocal peer learning with a social robot on children\u2019s learning and emotive engagement","volume":"150","author":"Chen","year":"2020","journal-title":"Computers & Education"},{"key":"10.1007\/s40593-025-00494-6_bib29","first-page":"258","article-title":"Reinforcement learning based feature selection for developing pedagogically effective tutorial dialogue tactics","author":"Chi","year":"2008","journal-title":"Proceedings of Educational Data Mining 2008 - 1st International Conference on Educational Data Mining"},{"key":"10.1007\/s40593-025-00494-6_bib30","series-title":"Lecture Notes in Computer Science \/ Information Systems and Applications, incl. Internet\/Web, and HCI: Vol. 6075. User Modeling, Adaptation, and Personalization: 18th International Conference, UMAP 2010, Big Island, HI, USA, June 20\u201324, 2010 ; proceedings","first-page":"147","article-title":"Inducing effective pedagogical strategies using learning context features","author":"Chi","year":"2010"},{"key":"10.1007\/s40593-025-00494-6_bib31","series-title":"Lecture Notes in Computer Science: Vol. 8474. Intelligent tutoring systems: 12th international conference, ITS 2014, Honolulu, HI, USA, June 5 - 9, 2014; proceedings","first-page":"210","article-title":"When is tutorial dialogue more effective than step-based tutoring?","author":"Chi","year":"2014"},{"issue":"1\u20132","key":"10.1007\/s40593-025-00494-6_bib32","doi-asserted-by":"crossref","first-page":"137","DOI":"10.1007\/s11257-010-9093-1","article-title":"Empirically evaluating the application of reinforcement learning to the induction of effective and adaptive pedagogical strategies","volume":"21","author":"Chi","year":"2011","journal-title":"User Modeling and User-Adapted Interaction"},{"issue":"2","key":"10.1007\/s40593-025-00494-6_bib33","first-page":"20","article-title":"Multi-armed bandits for intelligent tutoring systems","volume":"7","author":"Clement","year":"2015","journal-title":"Journal of Educational Data Mining"},{"key":"10.1007\/s40593-025-00494-6_bib34","series-title":"Statistical power analysis for the behavioral sciences","author":"Cohen","year":"1988"},{"key":"10.1007\/s40593-025-00494-6_bib35","first-page":"662","article-title":"A deep reinforcement learning approach to automatic formative feedback","author":"Condor","year":"2022","journal-title":"Proceedings of the 15th International Conference on Educational Data Mining"},{"key":"10.1007\/s40593-025-00494-6_bib36","article-title":"Towards an intelligent tutoring system for propositional proof construction","author":"Croy","year":"2008","journal-title":"Proceedings of the 2008 conference on Current Issues in Computing and Philosophy."},{"key":"10.1007\/s40593-025-00494-6_bib37","series-title":"2023 International Conference on Applied Intelligence and Sustainable Computing (ICAISC)","first-page":"1","article-title":"Research of Intelligent Tutoring System Based on Deep Learning Computer Technology","author":"Cui","year":"2023"},{"issue":"4","key":"10.1007\/s40593-025-00494-6_bib38","doi-asserted-by":"crossref","first-page":"812","DOI":"10.1213\/ANE.0000000000001596","article-title":"Publication Bias: The Elephant in the Review","volume":"123","author":"Dalton","year":"2016","journal-title":"Anesthesia and Analgesia"},{"issue":"2","key":"10.1007\/s40593-025-00494-6_bib39","doi-asserted-by":"crossref","first-page":"107","DOI":"10.3233\/DS-200028","article-title":"Reinforcement learning for personalization: A systematic literature review","volume":"3","author":"den Hengst","year":"2020","journal-title":"Data Science"},{"issue":"6","key":"10.1007\/s40593-025-00494-6_bib40","doi-asserted-by":"crossref","first-page":"2092","DOI":"10.1016\/j.eswa.2012.10.014","article-title":"Comparing strategies for modeling students learning styles through reinforcement learning in adaptive and intelligent educational systems: An experimental analysis","volume":"40","author":"Dor\u00e7a","year":"2013","journal-title":"Expert Systems with Applications"},{"issue":"4","key":"10.1007\/s40593-025-00494-6_bib41","doi-asserted-by":"crossref","first-page":"568","DOI":"10.1007\/s40593-019-00187-x","article-title":"Where\u2019s the reward?","volume":"29","author":"Doroudi","year":"2019","journal-title":"International Journal of Artificial Intelligence in Education"},{"issue":"9","key":"10.1007\/s40593-025-00494-6_bib42","doi-asserted-by":"crossref","first-page":"2419","DOI":"10.1007\/s10994-021-05961-4","article-title":"Challenges of real-world reinforcement learning: Definitions, benchmarks and analysis","volume":"110","author":"Dulac-Arnold","year":"2021","journal-title":"Machine Learning"},{"key":"10.1007\/s40593-025-00494-6_bib43","first-page":"388","article-title":"Zero-shot learning of hint policy via reinforcement learning and program synthesis","author":"Efremov","year":"2020","journal-title":"Proceedings of The 13th International Conference on Educational Data Mining (EDM 2020)"},{"issue":"7109","key":"10.1007\/s40593-025-00494-6_bib44","doi-asserted-by":"crossref","first-page":"629","DOI":"10.1136\/bmj.315.7109.629","article-title":"Bias in meta-analysis detected by a simple, graphical test","volume":"315","author":"Egger","year":"1997","journal-title":"BMJ : British Medical Journal"},{"issue":"1","key":"10.1007\/s40593-025-00494-6_bib45","doi-asserted-by":"crossref","first-page":"6637","DOI":"10.48084\/etasr.3905","article-title":"Design of an adaptive e-learning system based on multi-agent approach and reinforcement learning","volume":"11","author":"El Fazazi","year":"2021","journal-title":"Engineering, Technology & Applied Science Research"},{"issue":"3","key":"10.1007\/s40593-025-00494-6_bib46","doi-asserted-by":"crossref","first-page":"74","DOI":"10.3390\/informatics10030074","article-title":"Reinforcement learning in education: A literature review","volume":"10","author":"Fahad Mon","year":"2023","journal-title":"Informatics"},{"issue":"21","key":"10.1007\/s40593-025-00494-6_bib47","doi-asserted-by":"crossref","first-page":"23191","DOI":"10.1609\/aaai.v38i21.30365","article-title":"Online reinforcement learning-based pedagogical planning for narrative-centered learning environments","volume":"38","author":"Fahid","year":"2024","journal-title":"Proceedings of the AAAI Conference on Artificial Intelligence"},{"issue":"4","key":"10.1007\/s40593-025-00494-6_bib48","doi-asserted-by":"crossref","first-page":"e10271","DOI":"10.1371\/journal.pone.0010271","article-title":"Do pressures to publish increase scientists\u2019 bias? An empirical support from US States Data","volume":"5","author":"Fanelli","year":"2010","journal-title":"PLoS ONE"},{"key":"10.1007\/s40593-025-00494-6_bib49","series-title":"Icalt 2017: Ieee 17th International Conference on Advanced Learning Technologies: Proceedings: 3\u20137 July 2017, Timi\u015foara, Romania.","article-title":"Building adaptive tutoring model using artificial neural networks and reinforcement learning","author":"Fenza","year":"2017"},{"key":"10.1007\/s40593-025-00494-6_bib50","first-page":"605","article-title":"Transfer learning and representation discovery in intelligent tutoring systems","volume":"200","author":"Ferguson","year":"2009","journal-title":"Frontiers in Artificial Intelligence and Applications"},{"key":"10.1007\/s40593-025-00494-6_bib51","series-title":"2023 IEEE Frontiers in Education Conference (FIE)","first-page":"1","article-title":"Unleashing the Potential of Reinforcement Learning for Enhanced Personalized Education","author":"Fernandes","year":"2023"},{"key":"10.1007\/s40593-025-00494-6_bib52","series-title":"Modern educational paradigms for computer and engineering career: Proceedings: Edunine2019 - III IEEE World Engineering Education Conference : March 17 to 19, 2019, Lima, Peru","first-page":"1","article-title":"Proposal model for e-learning based on case based reasoning and reinforcement learning","author":"Flores","year":"2019"},{"key":"10.1007\/s40593-025-00494-6_bib53","series-title":"Programming and Software Engineering: Vol. 12149. Intelligent Tutoring Systems: 16th International Conference, ITS 2020, Athens, Greece, June 8\u201312, 2020, Proceedings (1st ed. 2020, Vol. 12149","first-page":"248","article-title":"An interactive recommender system based on reinforcement learning for improving emotional competences in educational groups","author":"Fotopoulou","year":"2020"},{"key":"10.1007\/s40593-025-00494-6_bib54","series-title":"ACM Digital Library, Proceedings of the 2nd International Conference on Computing and Wireless Communication Systems","first-page":"1","article-title":"Intelligent adapted e-learning system based on deep reinforcement learning","author":"El Fouki","year":"2017"},{"issue":"3\u20134","key":"10.1007\/s40593-025-00494-6_bib55","doi-asserted-by":"crossref","first-page":"219","DOI":"10.1561\/2200000071","article-title":"An introduction to deep reinforcement learning","volume":"11","author":"Fran\u00e7ois-Lavet","year":"2018","journal-title":"Foundations and Trends\u00ae in Machine Learning"},{"key":"10.1007\/s40593-025-00494-6_bib56","series-title":"Proceedings of the 14th Learning Analytics and Knowledge Conference","first-page":"426","article-title":"Finding Paths for Explainable MOOC Recommendation: A Learner Perspective","author":"Frej","year":"2024"},{"key":"10.1007\/s40593-025-00494-6_bib57","series-title":"Proceedings of the 2016 Conference on User Modeling Adaptation and Personalization","first-page":"131","article-title":"Adaptive training environment without prior knowledge","author":"Frenoy","year":"2016"},{"key":"10.1007\/s40593-025-00494-6_bib58","article-title":"Personalised human-robot co-adaptation in instructional settings using reinforcement learning","author":"Gao","year":"2017","journal-title":"IVA Workshop on Persuasive Embodied Agents for Behavior Change: PEACH 2017, August 27, Stockholm, Sweden."},{"key":"10.1007\/s40593-025-00494-6_bib59","series-title":"Ieee RO-MAN 2018: The 27th IEEE International Symposium on Robot and Human Interactive Communication","first-page":"705","article-title":"When robot personalisation does not help: Insights from a robot-supported learning study","author":"Gao","year":"2018"},{"key":"10.1007\/s40593-025-00494-6_bib60","first-page":"1504","article-title":"HOPE: Human-centric off-policy evaluation for e-learning and healthcare","author":"Gao","year":"2023","journal-title":"AAMAS \u201923, Proceedings of the 2023 International Conference on Autonomous Agents and Multiagent Systems"},{"key":"10.1007\/s40593-025-00494-6_bib61","doi-asserted-by":"crossref","unstructured":"Gao, H, Zeng, Y, Ma, B, Pan, Y. Improving knowledge learning through modelling students\u2019 practice-based cognitive processes, 2023. Advance online publication. 10.1007\/s12559-023-10201-z","DOI":"10.1007\/s12559-023-10201-z"},{"key":"10.1007\/s40593-025-00494-6_bib62","doi-asserted-by":"crossref","first-page":"737","DOI":"10.65109\/QVPV7447","article-title":"Using reinforcement learning to optimize the policies of an intelligent tutoring system for interpersonal skills training","author":"Georgila","year":"2019","journal-title":"AAMAS \u201919, Proceedings of the 18th International Conference on Autonomous Agents and MultiAgent Systems"},{"issue":"1","key":"10.1007\/s40593-025-00494-6_bib63","first-page":"1573","article-title":"RLPy: A value-function-based reinforcement learning framework for education and research","volume":"16","author":"Geramifard","year":"2015","journal-title":"Journal of Machine Learning Research"},{"key":"10.1007\/s40593-025-00494-6_bib64","first-page":"115","article-title":"Generating student feedback from time-series data using reinforcement learning","author":"Gkatzia","year":"2013","journal-title":"Proceedings of the 14th European Workshop on Natural Language Generation"},{"issue":"3","key":"10.1007\/s40593-025-00494-6_bib65","doi-asserted-by":"crossref","first-page":"1","DOI":"10.1145\/3580510","article-title":"Reinforced MOOCs concept recommendation in heterogeneous information networks","volume":"17","author":"Gong","year":"2023","journal-title":"ACM Transactions on the Web"},{"issue":"1","key":"10.1007\/s40593-025-00494-6_bib66","doi-asserted-by":"crossref","first-page":"3951","DOI":"10.1609\/aaai.v30i1.9914","article-title":"Affective personalization of a social robot tutor for children\u2019s second language skills","volume":"30","author":"Gordon","year":"2016","journal-title":"Proceedings of the AAAI Conference on Artificial Intelligence"},{"key":"10.1007\/s40593-025-00494-6_bib67","series-title":"APA handbooks in psychology. APA educational psychology handbook","first-page":"451","article-title":"Intelligent tutoring systems","author":"Graesser","year":"2012"},{"key":"10.1007\/s40593-025-00494-6_bib68","series-title":"Springer eBook Collection: Vol. 12677. Intelligent Tutoring Systems: 17th International Conference, ITS 2021, Virtual Event, June 7\u201311, 2021, Proceedings (1st ed. 2021, Vol. 12677","first-page":"439","article-title":"Towards smart edutainment applications for young children: A proposal","author":"Guran","year":"2021"},{"key":"10.1007\/s40593-025-00494-6_bib69","series-title":"Proceedings of the 2024 International Conference on Artificial Intelligence and Teacher Education","first-page":"117","article-title":"Genetic and Deep Reinforcement Learning-Based Intelligent Course Scheduling for Smart Education","author":"Haider","year":"2024"},{"issue":"3","key":"10.1007\/s40593-025-00494-6_bib70","doi-asserted-by":"crossref","first-page":"522","DOI":"10.1111\/bmsp.12199","article-title":"Curiosity-driven recommendation strategy for adaptive learning via deep reinforcement learning","volume":"73","author":"Han","year":"2020","journal-title":"The British Journal of Mathematical and Statistical Psychology"},{"key":"10.1007\/s40593-025-00494-6_bib71","series-title":"2023 IEEE International Conference on Systems, Man, and Cybernetics (SMC)","first-page":"1431","article-title":"Reinforcement Learning with Experience Sharing for Intelligent Educational Systems","author":"Hare","year":"2023"},{"issue":"3","key":"10.1007\/s40593-025-00494-6_bib72","doi-asserted-by":"crossref","first-page":"387","DOI":"10.1109\/TE.2024.3359001","article-title":"An Intelligent Serious Game for Digital Logic Education to Enhance Student Learning","volume":"67","author":"Hare","year":"2024","journal-title":"IEEE Transactions on Education"},{"key":"10.1007\/s40593-025-00494-6_bib73","first-page":"3207","article-title":"Deep reinforcement learning that matters","author":"Henderson","year":"2018","journal-title":"AAAI Conference On Artificial Intelligence (AAAI)"},{"key":"10.1007\/s40593-025-00494-6_bib74","first-page":"1248","article-title":"A self-organizing neuro-fuzzy q-network: Systematic design with offline hybrid learning","author":"Hostetter","year":"2023","journal-title":"AAMAS \u201923, Proceedings of the 2023 International Conference on Autonomous Agents and Multiagent Systems"},{"key":"10.1007\/s40593-025-00494-6_bib75","doi-asserted-by":"crossref","DOI":"10.1145\/3570945.3607301","article-title":"XAI to increase the effectiveness of an intelligent pedagogical agent","author":"Hostetter","year":"2023","journal-title":"ACM International Conference on Intelligent Virtual Agents (IVA2023)."},{"key":"10.1007\/s40593-025-00494-6_bib76","doi-asserted-by":"crossref","DOI":"10.1109\/FUZZ52849.2023.10309741","article-title":"Leveraging fuzzy logic towards more explainable reinforcement learning-induced pedagogical policies on intelligent tutoring systems","author":"Hostetter","year":"2023","journal-title":"IEEE International Conference on Fuzzy Systems (FUZZ 2023)."},{"key":"10.1007\/s40593-025-00494-6_bib77","series-title":"Dynamic programming and Markov processes","author":"Howard","year":"1960"},{"key":"10.1007\/s40593-025-00494-6_bib78","series-title":"ACM Digital Library, Cikm\u201919: Proceedings of the 28th ACM International Conference on Information & Knowledge Management","first-page":"1261","article-title":"Exploring multi-objective exercise recommendations in online education systems","author":"Huang","year":"2019"},{"key":"10.1007\/s40593-025-00494-6_bib79","series-title":"Otm 2003 Workshops: Otm Confederated International Workshops, HCI-SWWA, IPW, JTRES,WORM, WMS, and WRSM 2003, Catania, Sicily, Italy, November 3\u20137, 2003. Proceedings","article-title":"Navigating through the RLATES interface: A web-based adaptive and intelligent educational system","author":"Iglesias","year":"1995"},{"issue":"2","key":"10.1007\/s40593-025-00494-6_bib80","doi-asserted-by":"crossref","first-page":"223","DOI":"10.15388\/infedu.2003.17","article-title":"An experience applying reinforcement learning in a web-based adaptive and intelligent educational system","volume":"2","author":"Iglesias","year":"2003","journal-title":"Informatics in Education"},{"issue":"1","key":"10.1007\/s40593-025-00494-6_bib81","doi-asserted-by":"crossref","first-page":"89","DOI":"10.1007\/s10489-008-0115-1","article-title":"Learning teaching strategies in an adaptive and intelligent educational system through reinforcement learning","volume":"31","author":"Iglesias","year":"2009","journal-title":"Applied Intelligence"},{"issue":"4","key":"10.1007\/s40593-025-00494-6_bib82","doi-asserted-by":"crossref","first-page":"266","DOI":"10.1016\/j.knosys.2009.01.007","article-title":"Reinforcement learning of pedagogical policies in adaptive and intelligent educational systems","volume":"22","author":"Iglesias","year":"2009","journal-title":"Knowledge-Based Systems"},{"key":"10.1007\/s40593-025-00494-6_bib83","series-title":"The 6th Global Wireless Summit (GWS-2018): November 25\u201328, 2018, Mae Fah Luang University, Chiang Rai","first-page":"167","article-title":"Reinforcement learning for online learning recommendation system","author":"Intayoad","year":"2018"},{"issue":"4","key":"10.1007\/s40593-025-00494-6_bib84","doi-asserted-by":"crossref","first-page":"2917","DOI":"10.1007\/s11277-020-07199-0","article-title":"Reinforcement learning based on contextual bandits for personalized online learning recommendation systems","volume":"115","author":"Intayoad","year":"2020","journal-title":"Wireless Personal Communications"},{"key":"10.1007\/s40593-025-00494-6_bib85","article-title":"Reproducibility of benchmarked deep reinforcement learning tasks for continuous control","author":"Islam","year":"2017","journal-title":"ICML Reproducibility in Machine Learning Workshop, ICML\u201917."},{"key":"10.1007\/s40593-025-00494-6_bib86","doi-asserted-by":"crossref","first-page":"155123","DOI":"10.1109\/ACCESS.2021.3128578","article-title":"PAKES: A reinforcement learning-based personalized adaptability knowledge extraction strategy for adaptive learning systems","volume":"9","author":"Islam","year":"2021","journal-title":"IEEE Access"},{"key":"10.1007\/s40593-025-00494-6_bib87","unstructured":"JASP Team. (2021). JASP (Version 0.16.0.0) [Computer software]. https:\/\/jasp-stats.org\/"},{"key":"10.1007\/s40593-025-00494-6_bib88","series-title":"2019 IEEE Blocks and Beyond Workshop (B & B): B & B 2019 : October 18, 2019, Memphis, Tennessee, USA : Proceedings","first-page":"37","article-title":"It\u2019s not magic after all: Machine learning in Snap! using reinforcement learning","author":"Jatzlau","year":"2019"},{"key":"10.1007\/s40593-025-00494-6_bib89","series-title":"21st International Conference on Advances in ICT for Emerging Regions (ICTer) 2021: Conference proceedings : 02nd & 03rd of December 2021, University of Colombo, School of Computing, Colombo, Sri Lanka","first-page":"117","article-title":"English language trainer for non-native speakers using audio signal processing, reinforcement learning, and deep learning","author":"Jeewantha","year":"2021"},{"key":"10.1007\/s40593-025-00494-6_bib90","article-title":"Intelligent feedback polarity and timing selection in the Shufti intelligent tutoring system","author":"Johnson","year":"2013","journal-title":"International Conference on Computers in Education."},{"key":"10.1007\/s40593-025-00494-6_bib91","series-title":"Lecture Notes in Computer Science. ARTIFICIAL INTELLIGENCE IN EDUCATION: 22nd international conference, aied 2021","first-page":"215","article-title":"Evaluating critical reinforcement learning framework in the field","author":"Ju","year":"2021"},{"key":"10.1007\/s40593-025-00494-6_bib92","article-title":"More, May not the Better: Insights from Applying Deep Reinforcement Learning for Pedagogical Policy Induction","author":"Jung","year":"2024","journal-title":"Proceedings of the 17th International Conference on Educational Data Mining."},{"issue":"1\u20132","key":"10.1007\/s40593-025-00494-6_bib93","doi-asserted-by":"crossref","first-page":"99","DOI":"10.1016\/S0004-3702(98)00023-X","article-title":"Planning and acting in partially observable stochastic domains","volume":"101","author":"Kaelbling","year":"1998","journal-title":"Artificial Intelligence"},{"key":"10.1007\/s40593-025-00494-6_bib94","doi-asserted-by":"crossref","first-page":"1248","DOI":"10.1109\/TLT.2024.3372508","article-title":"Enhancing Medical Training Through Learning From Mistakes by Interacting With an Ill-Trained Reinforcement Learning Agent","volume":"17","author":"Kakdas","year":"2024","journal-title":"IEEE Transactions on Learning Technologies"},{"key":"10.1007\/s40593-025-00494-6_bib95","series-title":"2022 12th International Congress on Advanced Applied Informatics: Iiai-AAI 2022 : Kanazawa, Japan, 2\u20137 July 2022 : Proceedings","first-page":"230","article-title":"An analysis of educational cloud platforms using multi-agent learning","author":"Kandel","year":"2022"},{"key":"10.1007\/s40593-025-00494-6_bib96","series-title":"Springer eBook Collection: Vol. 12677. Intelligent Tutoring Systems: 17th International Conference, ITS 2021, Virtual Event, June 7\u201311, 2021, Proceedings (1st ed. 2021, Vol. 12677","first-page":"267","article-title":"Learning path construction using reinforcement learning and bloom\u2019s taxonomy","author":"Kim","year":"2021"},{"key":"10.1007\/s40593-025-00494-6_bib97","series-title":"Springer eBook Collection: Vol. 12164. Artificial Intelligence in Education: 21st International Conference, AIED 2020, Ifrane, Morocco, July 6\u201310, 2020, Proceedings, Part II (1st ed. 2020, Vol. 12164","first-page":"140","article-title":"Automated personalized feedback improves learning gains in an intelligent tutoring system","author":"Kochmar","year":"2020"},{"key":"10.1007\/s40593-025-00494-6_bib98","series-title":"Advances in Neural Information Processing Systems (Vol. 12).","article-title":"Actor-critic algorithms","author":"Konda","year":"1999"},{"issue":"3","key":"10.1007\/s40593-025-00494-6_bib99","first-page":"483","article-title":"Modelling an intelligent tutoring system using reinforcement learning","volume":"43","author":"Koroveshi","year":"2020","journal-title":"Knowledge - International Journal"},{"issue":"3","key":"10.1007\/s40593-025-00494-6_bib100","first-page":"10","article-title":"Training an intelligent tutoring system using reinforcement learning","volume":"19","author":"Koroveshi","year":"2021","journal-title":"International Journal of Computer Science and Information Security (IJCSIS)"},{"key":"10.1007\/s40593-025-00494-6_bib101","article-title":"RLTutor: Reinforcement learning based adaptive tutoring system by modeling virtual student with fewer interactions","author":"Kubotani","year":"2021","journal-title":"AI4EDU workshop at IJCAI2021."},{"key":"10.1007\/s40593-025-00494-6_bib102","series-title":"Springer eBooks Intelligent Technologies and Robotics: Vol. 989. Intelligent Communication, Control and Devices: Proceedings of ICICCD 2018 (1st ed. 2020, Vol. 989","first-page":"425","article-title":"An adaptive framework of learner model using learner characteristics for intelligent tutoring systems","author":"Kumar","year":"2020"},{"issue":"4","key":"10.1007\/s40593-025-00494-6_bib103","first-page":"1107","article-title":"Least-squares policy iteration","author":"Lagoudakis","year":"2003","journal-title":"The Journal of Machine Learning Research"},{"key":"10.1007\/s40593-025-00494-6_bib104","article-title":"Bandit algorithms","author":"Lattimore","year":"2020","journal-title":"Cambridge University Press"},{"key":"10.1007\/s40593-025-00494-6_bib105","doi-asserted-by":"crossref","first-page":"46","DOI":"10.1016\/j.jclinepi.2019.06.014","article-title":"Meta-analyses indexed in PsycINFO had a better completeness of reporting when they mention PRISMA","volume":"115","author":"Leclercq","year":"2019","journal-title":"Journal of Clinical Epidemiology"},{"key":"10.1007\/s40593-025-00494-6_bib106","unstructured":"Lee, J. in, & Brunskill, E. (2012). The impact on individualizing student models on necessary practice opportunities. In International Conference on Educational Data Mining (EDM), Chania, Greece."},{"key":"10.1007\/s40593-025-00494-6_bib107","series-title":"Proceedings \/ International Conference on Computers in Education: December 3 - 6, 2002, Aukland, New Zealand","first-page":"670","article-title":"A machine learning framework for an expert tutor construction","author":"Legaspi","year":"2002"},{"key":"10.1007\/s40593-025-00494-6_bib108","series-title":"ACM Digital Library, Lak22: 12th International Learning Analytics and Knowledge Conference","first-page":"294","article-title":"A novel video recommendation system for algebra: An effectiveness evaluation study","author":"Leite","year":"2022"},{"key":"10.1007\/s40593-025-00494-6_bib109","unstructured":"Lenhard, W., & Lenhard, A. (2017). Computation of effect sizes. Psychometrica. https:\/\/doi.org\/10.13140\/RG.2.2.17823.92329"},{"key":"10.1007\/s40593-025-00494-6_bib110","series-title":"Proceedings of the 30th ACM SIGKDD Conference on Knowledge Discovery and Data Mining","first-page":"1621","article-title":"Privileged Knowledge State Distillation for Reinforcement Learning-based Educational Path Recommendation","author":"Li","year":"2024"},{"issue":"9","key":"10.1007\/s40593-025-00494-6_bib111","first-page":"1823","article-title":"Learning Path Recommendation Based on Reinforcement Learning","volume":"32","author":"Li","year":"2024","journal-title":"Engineering Letters"},{"key":"10.1007\/s40593-025-00494-6_bib112","doi-asserted-by":"crossref","first-page":"16","DOI":"10.1016\/j.compedu.2019.01.003","article-title":"MOOC learners\u2019 demographics, self-regulated learning strategy, perceived learning and satisfaction: A structural equation modeling approach","volume":"132","author":"Li","year":"2019","journal-title":"Computers & Education"},{"issue":"2","key":"10.1007\/s40593-025-00494-6_bib113","doi-asserted-by":"crossref","first-page":"220","DOI":"10.3102\/10769986221129847","article-title":"Deep reinforcement learning for adaptive learning systems","volume":"48","author":"Li","year":"2023","journal-title":"Journal of Educational and Behavioral Statistics"},{"issue":"34","key":"10.1007\/s40593-025-00494-6_bib114","doi-asserted-by":"crossref","first-page":"24369","DOI":"10.1007\/s00521-023-08989-w","article-title":"Sim-GAIL: A generative adversarial imitation learning approach of student modelling for intelligent tutoring systems","volume":"35","author":"Li","year":"2023","journal-title":"Neural Computing and Applications"},{"key":"10.1007\/s40593-025-00494-6_bib115","series-title":"2024 4th International Signal Processing, Communications and Engineering Management Conference (ISPCEM)","first-page":"801","article-title":"Research on personalized learning path recommendation model of artificial intelligence in new business","author":"Liang","year":"2024"},{"key":"10.1007\/s40593-025-00494-6_bib116","doi-asserted-by":"crossref","first-page":"120757","DOI":"10.1109\/ACCESS.2020.3006254","article-title":"A review on interactive reinforcement learning from human social feedback","volume":"8","author":"Lin","year":"2020","journal-title":"IEEE Access"},{"key":"10.1007\/s40593-025-00494-6_bib117","series-title":"Proceedings of 2018 IEEE International Conference on Teaching, Assessment, and Learning for Engineering (TALE2018)","first-page":"1079","article-title":"Towards smart educational recommendations with reinforcement learning in classroom","author":"Liu","year":"2018"},{"key":"10.1007\/s40593-025-00494-6_bib118","series-title":"ACM Digital Library, Proceedings of the 25th ACM SIGKDD International Conference on Knowledge Discovery & Data Mining","first-page":"627","article-title":"Exploiting cognitive structure for adaptive learning","author":"Liu","year":"2019"},{"key":"10.1007\/s40593-025-00494-6_bib119","series-title":"2nd International Conference on Mobile Networks and Wireless Communications (ICMNWC-2022)","first-page":"1","article-title":"The dynamic mode construction of mixed english learning based on reinforcement learning","author":"Liu","year":"2022"},{"key":"10.1007\/s40593-025-00494-6_bib120","series-title":"EDULEARN Proceedings, EDULEARN23 Proceedings","first-page":"7148","article-title":"Enhancing STEM Education using Machine Learning and Reinforcement Learning Techniques for Educational Software and Serious Games","author":"Liu","year":"2023"},{"key":"10.1007\/s40593-025-00494-6_bib121","series-title":"ACM Digital Library, Proceedings of the 29th ACM SIGKDD Conference on Knowledge Discovery and Data Mining","first-page":"1441","article-title":"Meta multi-agent exercise recommendation: A game application perspective","author":"Liu","year":"2023"},{"key":"10.1007\/s40593-025-00494-6_bib122","series-title":"ACM Digital Library, Lak21: 11th International Learning Analytics and Knowledge Conference","first-page":"397","article-title":"Recommendation for effective standardized exam preparation","author":"Loh","year":"2021"},{"issue":"3","key":"10.1007\/s40593-025-00494-6_bib123","doi-asserted-by":"crossref","DOI":"10.1002\/jrsm.1310","article-title":"Dealing with effect size multiplicity in systematic reviews and meta-analyses","volume":"9","author":"L\u00f3pez-L\u00f3pez","year":"2018","journal-title":"Research Synthesis Methods"},{"key":"10.1007\/s40593-025-00494-6_bib124","article-title":"Learning expert models for educationally relevant tasks using reinforcement learning","author":"Maclellan","year":"2021","journal-title":"International Conference on Educational Data Mining (EDM)."},{"issue":"10","key":"10.1007\/s40593-025-00494-6_bib125","doi-asserted-by":"crossref","first-page":"3921","DOI":"10.1007\/s12652-019-01627-1","article-title":"Finding optimal pedagogical content in an adaptive e-learning platform using a new recommendation approach and reinforcement learning","volume":"11","author":"Madani","year":"2020","journal-title":"Journal of Ambient Intelligence and Humanized Computing"},{"issue":"3","key":"10.1007\/s40593-025-00494-6_bib126","doi-asserted-by":"crossref","DOI":"10.1002\/cl2.1256","article-title":"Using selection models to assess sensitivity to publication bias: A tutorial and call for more routine use","volume":"18","author":"Maier","year":"2022","journal-title":"Campbell Systematic Reviews"},{"key":"10.1007\/s40593-025-00494-6_bib127","article-title":"Personalized intelligent tutoring system using reinforcement learning","author":"Malpani","year":"2011","journal-title":"Proceedings of the Twenty-Fourth International Florida Artificial Intelligence Research Society Conference, May 18\u201320, 2011, Palm Beach, Florida, USA."},{"key":"10.1007\/s40593-025-00494-6_bib128","series-title":"AAMAS \u201914, Proceedings of the 2014 International Conference on Autonomous Agents and Multi-Agent Systems","first-page":"1077","article-title":"Offline policy evaluation across representations with applications to educational games","author":"Mandel","year":"2014"},{"key":"10.1007\/s40593-025-00494-6_bib129","series-title":"Better education through improved reinforcement learning [Doctoral dissertation","author":"Mandel","year":"2017"},{"key":"10.1007\/s40593-025-00494-6_bib130","series-title":"Lecture Notes in Computer Science: Vol. 3220. Intelligent Tutoring Systems: 7th International Conference, ITS 2004 Proceedings","first-page":"564","article-title":"Agentx: Using reinforcement learning to improve the effectiveness of intelligent tutoring systems","author":"Martin","year":"2004"},{"issue":"2","key":"10.1007\/s40593-025-00494-6_bib131","doi-asserted-by":"crossref","first-page":"133","DOI":"10.1177\/0962280211432219","article-title":"A practical introduction to multivariate meta-analysis","volume":"22","author":"Mavridis","year":"2013","journal-title":"Statistical Methods in Medical Research"},{"issue":"8","key":"10.1007\/s40593-025-00494-6_bib132","doi-asserted-by":"crossref","first-page":"9325","DOI":"10.1007\/s10639-022-11129-x","article-title":"Pilot study of an intervention based on an intelligent tutoring system (ITS) for instructing mathematical skills of students with ASD and\/or ID","volume":"28","author":"Mazon","year":"2023","journal-title":"Education and Information Technologies"},{"key":"10.1007\/s40593-025-00494-6_bib133","doi-asserted-by":"crossref","DOI":"10.1016\/j.caeo.2024.100175","article-title":"A scoping review of reinforcement learning in education","volume":"6","author":"Memarian","year":"2024","journal-title":"Computers and Education Open"},{"key":"10.1007\/s40593-025-00494-6_bib134","series-title":"2010 Second International Workshop on Education Technology and Computer Science: Etcs 2010 ; Wuhan, China, 6 - 7 March 2010 ; [proceedings","first-page":"288","article-title":"Course-scheduling algorithm of option-based hierarchical reinforcement learning","author":"Ming","year":"2010"},{"issue":"5","key":"10.1007\/s40593-025-00494-6_bib135","doi-asserted-by":"crossref","first-page":"6389","DOI":"10.1007\/s11042-021-11806-y","article-title":"Ralf: An adaptive reinforcement learning framework for teaching dyslexic students","volume":"81","author":"Minoofam","year":"2022","journal-title":"Multimedia Tools and Applications"},{"key":"10.1007\/s40593-025-00494-6_bib136","series-title":"2009 IEEE\/WIC\/ACM International Joint Conferences on Web Intelligence (WI) and Intelligent Agent Technologies (IAT): Wi-IAT 2009","first-page":"215","article-title":"Adaptive learning based on exercises fitness degree","author":"Mirea","year":"2009"},{"key":"10.1007\/s40593-025-00494-6_bib137","series-title":"Human and environment friendly robots with high intelligence and emotional quotients: Proceedings","first-page":"1420","article-title":"Active learning from cross perceptual aliasing caused by direct teaching","author":"Mishima","year":"1999"},{"key":"10.1007\/s40593-025-00494-6_bib138","series-title":"Lecture notes in computer science Lecture notes in artificial intelligence: Vol. 7926. Artificial intelligence in education: 16th international conference, AIED 2013 proceedings","first-page":"828","article-title":"A markov decision process model of tutorial intervention in task-oriented dialogue","author":"Mitchell","year":"2013"},{"issue":"7540","key":"10.1007\/s40593-025-00494-6_bib139","doi-asserted-by":"crossref","first-page":"529","DOI":"10.1038\/nature14236","article-title":"Human-level control through deep reinforcement learning","volume":"518","author":"Mnih","year":"2015","journal-title":"Nature"},{"issue":"1","key":"10.1007\/s40593-025-00494-6_bib140","doi-asserted-by":"crossref","first-page":"1","DOI":"10.1561\/2200000086","article-title":"Model-based reinforcement learning: A survey","volume":"16","author":"Moerland","year":"2023","journal-title":"Foundations and Trends\u00ae in Machine Learning"},{"key":"10.1007\/s40593-025-00494-6_bib141","series-title":"2024 International Conference on Artificial Intelligence and Quantum Computation-Based Sensor Application (ICAIQSA)","first-page":"1","article-title":"Increasing Learner Engagement in English Language Acquisition Through AI-Powered Gamification","author":"Mohana","year":"2024"},{"issue":"1","key":"10.1007\/s40593-025-00494-6_bib142","doi-asserted-by":"crossref","first-page":"1","DOI":"10.1287\/mnsc.28.1.1","article-title":"State of the Art\u2014A Survey of Partially Observable Markov Decision Processes: Theory, Models, and Algorithms","volume":"28","author":"Monahan","year":"1982","journal-title":"Management Science"},{"key":"10.1007\/s40593-025-00494-6_bib143","series-title":"ACM Proceedings of the Fifth Annual ACM Conference on Learning at Scale","first-page":"1","article-title":"Combining adaptivity with progression ordering for intelligent tutoring systems","author":"Mu","year":"2018"},{"key":"10.1007\/s40593-025-00494-6_bib144","series-title":"Springer eBook Collection: Vol. 12677. Intelligent Tutoring Systems: 17th International Conference, ITS 2021, Virtual Event, June 7\u201311, 2021, Proceedings (1st ed. 2021, Vol. 12677","first-page":"430","article-title":"Automatic adaptive sequencing in a webgame","author":"Mu","year":"2021"},{"key":"10.1007\/s40593-025-00494-6_bib145","unstructured":"Murphy, K. (2025). Reinforcement learning: An overview. http:\/\/arxiv.org\/pdf\/2412.05265"},{"issue":"5","key":"10.1007\/s40593-025-00494-6_bib146","doi-asserted-by":"crossref","DOI":"10.14569\/IJACSA.2023.0140528","article-title":"Towards an adaptive e-learning system based on deep learner profile, machine learning approach, and reinforcement learning","volume":"14","author":"Mustapha","year":"2023","journal-title":"International Journal of Advanced Computer Science and Applications"},{"issue":"3","key":"10.1007\/s40593-025-00494-6_bib147","doi-asserted-by":"crossref","first-page":"221","DOI":"10.1080\/00220671.2017.1289775","article-title":"Integrated STEM defined: Contexts, challenges, and the future","volume":"110","author":"Nadelson","year":"2017","journal-title":"The Journal of Educational Research"},{"key":"10.1007\/s40593-025-00494-6_bib148","series-title":"2019 Third IEEE International Conference on Robotic Computing (IRC)","first-page":"590","article-title":"Review of deep reinforcement learning for robot manipulation","author":"Nguyen","year":"2019"},{"key":"10.1007\/s40593-025-00494-6_bib149","first-page":"688","article-title":"Understanding the impact of reinforcement learning personalization on subgroups of students in math tutoring","author":"Nie","year":"2023"},{"key":"10.1007\/s40593-025-00494-6_bib150","series-title":"2019 IEEE 14th International Conference on Industrial and Information Systems: (ICIIS) : 18th-20th December, 2019 : Conference proceedings","first-page":"446","article-title":"Athwel: Gamification supportive tool for special educational centers in Sri Lanka","author":"Nisansala","year":"2019"},{"key":"10.1007\/s40593-025-00494-6_bib151","series-title":"2022 IEEE 25th International Conference on Computer Supported Cooperative Work in Design (CSCWD)","first-page":"299","article-title":"Get a sense of accomplishment in doing exercises: A reinforcement learning perspective","author":"Niu","year":"2022"},{"issue":"4","key":"10.1007\/s40593-025-00494-6_bib152","doi-asserted-by":"crossref","DOI":"10.1007\/BF00168958","article-title":"Intelligent tutoring systems: an overview","volume":"4","author":"Nwana","year":"1990","journal-title":"Artificial Intelligence Review"},{"key":"10.1007\/s40593-025-00494-6_bib153","series-title":"Hri \u201822: Proceedings of the 2022 ACM\/IEEE International Conference on Human-Robot Interaction","first-page":"963","article-title":"K-Qbot: Language learning chatbot based on reinforcement learning","author":"Oralbayeva","year":"2022"},{"key":"10.1007\/s40593-025-00494-6_bib154","series-title":"Lecture Notes in Computer Science: Vol. 13891. Augmented Intelligence and Intelligent Tutoring Systems: 19th International Conference, ITS 2023, Corfu, Greece, June 2\u20135, 2023, Proceedings (1st ed. 2023, Vol. 13891","first-page":"16","article-title":"Recommending mathematical tasks based on reinforcement learning and item response theory","author":"Orsoni","year":"2023"},{"issue":"2","key":"10.1007\/s40593-025-00494-6_bib155","doi-asserted-by":"crossref","first-page":"55","DOI":"10.32591\/coas.ojit.0402.03055o","article-title":"Reinforcement learning approach for adaptive e-learning based on multiple learner characteristics","volume":"4","author":"Oyuga Anne","year":"2021","journal-title":"Open Journal for Information Technology"},{"key":"10.1007\/s40593-025-00494-6_bib156","article-title":"Prisma 2020 explanation and elaboration: Updated guidance and exemplars for reporting systematic reviews","volume":"372","author":"Page","year":"2021","journal-title":"BMJ (Clinical Research Ed.)"},{"key":"10.1007\/s40593-025-00494-6_bib157","series-title":"2024 2nd International Conference on Mechatronics, IoT and Industrial Informatics (ICMIII)","first-page":"599","article-title":"Application of Reinforcement Learning Algorithm in Personalized Music Teaching","author":"Pan","year":"2024"},{"issue":"12","key":"10.1007\/s40593-025-00494-6_bib158","doi-asserted-by":"crossref","DOI":"10.1371\/journal.pone.0083138","article-title":"Evaluation of the endorsement of the preferred reporting items for systematic reviews and meta-analysis (PRISMA) statement on the quality of published systematic review and meta-analyses","volume":"8","author":"Panic","year":"2013","journal-title":"PLoS ONE"},{"issue":"3","key":"10.1007\/s40593-025-00494-6_bib159","doi-asserted-by":"crossref","first-page":"441","DOI":"10.1287\/moor.12.3.441","article-title":"The Complexity of Markov Decision Processes","volume":"12","author":"Papadimitriou","year":"1987","journal-title":"Mathematics of Operations Research"},{"key":"10.1007\/s40593-025-00494-6_bib160","doi-asserted-by":"crossref","first-page":"687","DOI":"10.1609\/aaai.v33i01.3301687","article-title":"A model-free affective reinforcement learning approach to personalization of an autonomous social robot companion for early literacy education","volume":"33","author":"Park","year":"2019","journal-title":"Proceedings of the AAAI Conference on Artificial Intelligence"},{"key":"10.1007\/s40593-025-00494-6_bib161","series-title":"2021 International Conference on Computing, Communication and Green Engineering (CCGE2021)","first-page":"1","article-title":"Application for multi-agent system: A case of customised elearning","author":"Patel","year":"2021"},{"key":"10.1007\/s40593-025-00494-6_bib162","series-title":"2022 XVII Latin American Conference on Learning Technologies (LACLO)","first-page":"1","article-title":"Reinforcement learning for estimating student proficiency in math word problems","author":"Perez","year":"2022"},{"issue":"16","key":"10.1007\/s40593-025-00494-6_bib163","doi-asserted-by":"crossref","first-page":"21015","DOI":"10.1007\/s10639-024-12699-8","article-title":"Emotions as implicit feedback for adapting difficulty in tutoring systems based on reinforcement learning","volume":"29","author":"P\u00e9rez","year":"2024","journal-title":"Education and Information Technologies"},{"key":"10.1007\/s40593-025-00494-6_bib164","first-page":"1","article-title":"Optimization of a tutoring system from a fixed set of data","author":"Pietquin","year":"2011","journal-title":"SLaTE 2011"},{"key":"10.1007\/s40593-025-00494-6_bib165","series-title":"Lecture Notes in Computer Science: Vol. 14798. Generative Intelligence and Intelligent Tutoring Systems","doi-asserted-by":"crossref","first-page":"117","DOI":"10.1007\/978-3-031-63028-6_10","article-title":"Individualised Mathematical Task Recommendations Through Intended Learning Outcomes and Reinforcement Learning","author":"P\u00f6gelt","year":"2024"},{"key":"10.1007\/s40593-025-00494-6_bib166","series-title":"2012 International Conference on Computer Communication and Informatics (ICCCI 2012)","first-page":"1","article-title":"Learning agent based knowledge management in intelligent tutoring system","author":"Priya","year":"2012"},{"key":"10.1007\/s40593-025-00494-6_bib167","series-title":"2020 IEEE International Conference on Big Data (Big Data)","first-page":"5201","article-title":"A deep reinforcement learning framework for instructional sequencing","author":"Pu","year":"2020"},{"key":"10.1007\/s40593-025-00494-6_bib168","series-title":"Markov decision processes: Discrete stochastic dynamic programming. Wiley series in probability and mathematical statistics. Applied probability and statistics section.","author":"Puterman","year":"2005"},{"key":"10.1007\/s40593-025-00494-6_bib169","series-title":"2014 IEEE International Conference on MOOC, Innovation and Technology in Education (MITE)","first-page":"285","article-title":"Reinforcement learning approach towards effective content recommendation in MOOC environments","author":"Raghuveer","year":"2014"},{"key":"10.1007\/s40593-025-00494-6_bib170","article-title":"Adapting difficulty levels in personalized robot-child tutoring interactions","author":"Ramachandran","year":"2014","journal-title":"Workshops at the Twenty-Eighth AAAI Conference on Artificial Intelligence"},{"key":"10.1007\/s40593-025-00494-6_bib171","series-title":"2021 30th IEEE International Conference on Robot & Human Interactive Communication (RO-MAN)","first-page":"250","article-title":"Effects of an adaptive robot encouraging teamwork on students\u2019 learning","author":"Ravari","year":"2021"},{"key":"10.1007\/s40593-025-00494-6_bib172","first-page":"5","article-title":"Accelerating human learning with deep reinforcement learning","author":"Reddy","year":"2017","journal-title":"NIPS\u201917 Workshop: Teaching Machines, Robots, and Humans"},{"key":"10.1007\/s40593-025-00494-6_bib173","series-title":"ACM Digital Library, Proceedings of the 23rd ACM International Conference on Intelligent Virtual Agents","first-page":"1","article-title":"Towards an Adaptive Pedagogical Agent in a Reading Intervention Using Reinforcement Learning","author":"Riedmann","year":"2023"},{"key":"10.1007\/s40593-025-00494-6_bib174","series-title":"Lecture Notes in Artificial Intelligence: Vol. 14992. Ki 2024: Advances in Artificial Intelligence: 47th German Conference on AI, W\u00fcrzburg, Germany, September 25\u201327, 2024, Proceedings (1st ed. 2024, Vol. 1410).","article-title":"Uli-RL: A Real-World Deep Reinforcement Learning Pedagogical Agent for Children","author":"Riedmann","year":"2024"},{"issue":"4","key":"10.1007\/s40593-025-00494-6_bib175","doi-asserted-by":"crossref","first-page":"789","DOI":"10.1111\/j.1467-985X.2008.00593.x","article-title":"Multivariate meta-analysis: The effect of ignoring within-study correlation","volume":"172","author":"Riley","year":"2009","journal-title":"Journal of the Royal Statistical Society Series a: Statistics in Society"},{"key":"10.1007\/s40593-025-00494-6_bib176","first-page":"503","article-title":"Bayesian inverse reinforcement learning for modeling conversational agents in a virtual environment","author":"Rojas-Barahona","year":"2014"},{"key":"10.1007\/s40593-025-00494-6_bib177","series-title":"Lecture notes in computer science Lecture notes in artificial intelligence: Vol. 9112. Artificial intelligence in education: 17th international conference, AIED 2015 proceedings","first-page":"419","article-title":"Improving student problem solving in narrative-centered learning environments: A modular reinforcement learning framework","author":"Rowe","year":"2015"},{"key":"10.1007\/s40593-025-00494-6_bib178","series-title":"Ieee RO-MAN 2018: The 27th IEEE International Symposium on Robot and Human Interactive Communication","first-page":"294","article-title":"A reinforcement learning model for robots as teachers","author":"Roy","year":"2018"},{"issue":"5","key":"10.1007\/s40593-025-00494-6_bib179","doi-asserted-by":"crossref","first-page":"3023","DOI":"10.1007\/s10994-023-06423-9","article-title":"Reinforcement learning tutor better supported lower performers in a math task","volume":"113","author":"Ruan","year":"2024","journal-title":"Machine Learning"},{"key":"10.1007\/s40593-025-00494-6_bib180","doi-asserted-by":"crossref","DOI":"10.1016\/j.compedu.2021.104426","article-title":"Large scale analytics of global and regional MOOC providers: Differences in learners\u2019 demographics, preferences, and perceptions","volume":"180","author":"Ruip\u00e9rez-Valiente","year":"2022","journal-title":"Computers & Education"},{"key":"10.1007\/s40593-025-00494-6_bib181","series-title":"IFIP International Federation for Information Processing: Vol. 241. Home informatics and telematics: Ict for the next billion: Proceeding of IFIP TC 9, WG 9.3 HOIT 2007 Conference","first-page":"65","article-title":"Intelligent tutoring systems using reinforcement learning to teach autistic students","author":"Sarma","year":"2007"},{"key":"10.1007\/s40593-025-00494-6_bib182","series-title":"Lecture Notes in Computer Science: Vol. 10331. Artificial Intelligence in Education: 18th International Conference, AIED 2017 Proceedings","first-page":"323","article-title":"Balancing learning and engagement in game-based learning environments with multi-objective reinforcement learning","author":"Sawyer","year":"2017"},{"key":"10.1007\/s40593-025-00494-6_bib183","series-title":"Lecture Notes in Computer Science: Vol. 14829. Artificial Intelligence in Education","doi-asserted-by":"crossref","first-page":"280","DOI":"10.1007\/978-3-031-64302-6_20","article-title":"Improving the Validity of Automatically Generated Feedback via Reinforcement Learning","author":"Scarlatos","year":"2024"},{"key":"10.1007\/s40593-025-00494-6_bib184","series-title":"Lecture Notes in Computer Science: Vol. 14200. Responsive and Sustainable Educational Futures: 18th European Conference on Technology Enhanced Learning, EC-TEL 2023 Proceedings (1st ed. 2023, Vol. 14200","first-page":"383","article-title":"Learning to give useful hints: Assistance action evaluation and policy improvements","author":"Schmucker","year":"2023"},{"key":"10.1007\/s40593-025-00494-6_bib185","series-title":"Proceedings of Machine Learning Research, Proceedings of the 32nd International Conference on Machine Learning","first-page":"1889","article-title":"Trust region policy optimization","author":"Schulman","year":"2015"},{"key":"10.1007\/s40593-025-00494-6_bib186","unstructured":"Schulman, J., Wolski, F., Dhariwal, P., Radford, A., & Klimov, O. (2017). Proximal policy optimization algorithms. https:\/\/arxiv.org\/pdf\/1707.06347.pdf"},{"key":"10.1007\/s40593-025-00494-6_bib187","doi-asserted-by":"crossref","DOI":"10.1016\/j.eswa.2023.120495","article-title":"Reinforcement learning algorithms: A brief survey","volume":"231","author":"Shakya","year":"2023","journal-title":"Expert Systems with Applications"},{"key":"10.1007\/s40593-025-00494-6_bib188","article-title":"Designing Simulated Students to Emulate Learner Activity Data in an Open-Ended Learning Environment","author":"Sharma","year":"2024","journal-title":"Proceedings of the 17th International Conference on Educational Data Mining."},{"key":"10.1007\/s40593-025-00494-6_bib189","series-title":"Advances in Intelligent Systems and Computing: Vol. 723. The International Conference on Advanced Machine Learning Technologies and Applications (AMLTA2018)","doi-asserted-by":"crossref","first-page":"221","DOI":"10.1007\/978-3-319-74690-6_22","article-title":"A reinforcement learning-based adaptive learning system","author":"Shawky","year":"2018"},{"key":"10.1007\/s40593-025-00494-6_bib190","series-title":"Studies in Computational Intelligence: Volume 801. Machine Learning Paradigms: Theory and Application","first-page":"169","article-title":"Towards a personalized learning experience using reinforcement learning","author":"Shawky","year":"2019"},{"key":"10.1007\/s40593-025-00494-6_bib191","series-title":"ACM Conferences, Proceedings of the 26th Conference on User Modeling, Adaptation and Personalization","first-page":"43","article-title":"Improving learning & reducing time: A constrained action-based reinforcement learning approach","author":"Shen","year":"2018"},{"key":"10.1007\/s40593-025-00494-6_bib192","series-title":"Vol. 10948. Artificial Intelligence in Education: 19th International Conference, AIED 2018 Proceedings, Part II","first-page":"327","article-title":"Empirically evaluating the effectiveness of pomdp vs. mdp towards the pedagogical strategies induction","author":"Shen","year":"2018"},{"key":"10.1007\/s40593-025-00494-6_bib193","series-title":"2021 International Conference on Signal Processing and Machine Learning: Conf-SPML 2021 Proceedings","first-page":"186","article-title":"Using q-learning to personalize pedagogical policies for addition problems","author":"Shen","year":"2021"},{"issue":"3","key":"10.1007\/s40593-025-00494-6_bib194","first-page":"27","article-title":"Exploring induced pedagogical strategies through a markov decision process framework: Lessons learned","volume":"10","author":"Shen","year":"2018","journal-title":"Journal of Educational Data Mining"},{"issue":"1","key":"10.1007\/s40593-025-00494-6_bib195","doi-asserted-by":"crossref","first-page":"216","DOI":"10.3758\/s13428-021-01602-9","article-title":"Building an intelligent recommendation system for personalized test scheduling in computerized assessments: A reinforcement learning approach","volume":"54","author":"Shin","year":"2022","journal-title":"Behavior Research Methods"},{"key":"10.1007\/s40593-025-00494-6_bib196","doi-asserted-by":"crossref","first-page":"24","DOI":"10.1186\/s13643-015-0004-8","article-title":"Quantifying the risk of error when interpreting funnel plots","volume":"4","author":"Simmonds","year":"2015","journal-title":"Systematic Reviews"},{"key":"10.1007\/s40593-025-00494-6_bib197","article-title":"Reinforcement learning for education: Opportunities and challenges [Workshop]","author":"Singla","year":"2021","journal-title":"International Conference on Educational Data Mining (EDM)."},{"issue":"3","key":"10.1007\/s40593-025-00494-6_bib198","doi-asserted-by":"crossref","first-page":"56","DOI":"10.5815\/ijmecs.2024.03.05","article-title":"Automatic Real-Time Adaptation of Training Session Difficulty Using Rules and Reinforcement Learning in the AI-VT ITS","volume":"16","author":"Soto Forero","year":"2024","journal-title":"International Journal of Modern Education and Computer Science"},{"key":"10.1007\/s40593-025-00494-6_bib199","article-title":"A reinforcement learning approach to adaptive remediation in online training","author":"Spain","year":"2021","journal-title":"The Journal of Defense Modeling and Simulation: Applications, Methodology, Technology"},{"key":"10.1007\/s40593-025-00494-6_bib200","first-page":"71","article-title":"The hint factory: Automatic generation of contextualized help for existing computer aided instruction","author":"Stamper","year":"2008","journal-title":"Proceedings of the 9th International Conference on Intelligent Tutoring Systems Young Researchers Track"},{"issue":"5","key":"10.1007\/s40593-025-00494-6_bib201","doi-asserted-by":"crossref","first-page":"581","DOI":"10.1177\/1948550617693062","article-title":"Limitations of PET-PEESE and Other Meta-Analysis Methods","volume":"8","author":"Stanley","year":"2017","journal-title":"Social Psychological and Personality Science"},{"issue":"1","key":"10.1007\/s40593-025-00494-6_bib202","doi-asserted-by":"crossref","first-page":"60","DOI":"10.1002\/jrsm.1095","article-title":"Meta-regression approximations to reduce publication selection bias","volume":"5","author":"Stanley","year":"2014","journal-title":"Research Synthesis Methods"},{"issue":"7304","key":"10.1007\/s40593-025-00494-6_bib203","doi-asserted-by":"crossref","first-page":"101","DOI":"10.1136\/bmj.323.7304.101","article-title":"Systematic reviews in health care: Investigating and dealing with publication and other biases in meta-analysis","volume":"323","author":"Sterne","year":"2001","journal-title":"BMJ : British Medical Journal"},{"key":"10.1007\/s40593-025-00494-6_bib204","doi-asserted-by":"crossref","DOI":"10.1136\/bmj.d4002","article-title":"Recommendations for examining and interpreting funnel plot asymmetry in meta-analyses of randomised controlled trials","volume":"343","author":"Sterne","year":"2011","journal-title":"BMJ (Clinical Research Ed.)"},{"key":"10.1007\/s40593-025-00494-6_bib205","series-title":"Icassp 2013 - 2013 IEEE International Conference on Acoustics, Speech and Signal Processing","first-page":"8213","article-title":"A dialogue game framework with personalized training using reinforcement learning for computer-assisted language learning","author":"Su","year":"2013"},{"key":"10.1007\/s40593-025-00494-6_bib206","series-title":"2024 International Conference on Telecommunications and Power Electronics (TELEPE)","first-page":"579","article-title":"Simulation and Optimization of Physical Education Teaching Based on Virtual Reality Technology and Reinforcement Learning Algorithms","author":"Sun","year":"2024"},{"key":"10.1007\/s40593-025-00494-6_bib207","series-title":"Reinforcement learning: An introduction","author":"Sutton","year":"2018"},{"key":"10.1007\/s40593-025-00494-6_bib208","series-title":"Fie 2022 Proceedings","first-page":"1","article-title":"Classroom evaluation of a gamified adaptive tutoring system","author":"Tang","year":"2022"},{"issue":"1","key":"10.1007\/s40593-025-00494-6_bib209","doi-asserted-by":"crossref","first-page":"108","DOI":"10.1111\/bmsp.12144","article-title":"A reinforcement learning approach to personalized learning recommendation systems","volume":"72","author":"Tang","year":"2019","journal-title":"The British Journal of Mathematical and Statistical Psychology"},{"issue":"1","key":"10.1007\/s40593-025-00494-6_bib210","doi-asserted-by":"crossref","DOI":"10.5041\/RMMJ.10518","article-title":"Against Over-reliance on PRISMA Guidelines for Meta-analytical Studies","volume":"15","author":"Teixeira da Silva","year":"2024","journal-title":"Rambam Maimonides Medical Journal"},{"key":"10.1007\/s40593-025-00494-6_bib211","series-title":"Proceedings of the main conference on Human Language Technology Conference of the North American Chapter of the Association of Computational Linguistics -","first-page":"272","article-title":"Comparing the utility of state features in spoken dialogue using reinforcement learning","author":"Tetreault","year":"2006"},{"key":"10.1007\/s40593-025-00494-6_bib212","article-title":"Using reinforcement learning to build a better model of dialogue state","author":"Tetreault","year":"2006","journal-title":"EACL 2006, 11st Conference of the European Chapter of the Association for Computational Linguistics Proceedings."},{"issue":"1","key":"10.1007\/s40593-025-00494-6_bib213","first-page":"107","article-title":"Using reinforcement learning to introduce artificial intelligence in the CS curriculum","volume":"18","author":"Thede","year":"2002","journal-title":"Journal of Computing Sciences in Colleges"},{"key":"10.1007\/s40593-025-00494-6_bib214","series-title":"Technology and the curriculum: Summer 2019.","article-title":"Enhancing STEM learning through technology","author":"Vahidy","year":"2019"},{"key":"10.1007\/s40593-025-00494-6_bib215","first-page":"17","article-title":"Developing pedagogically effective tutorial dialogue tactics: Experiments and a testbed","author":"VanLehn","year":"2007","journal-title":"SLaTE-2007"},{"key":"10.1007\/s40593-025-00494-6_bib216","article-title":"Towards Scalable Adaptive Learning with Graph Neural Networks and Reinforcement Learning","author":"Vassoyan","year":"2023","journal-title":"Proceedings of the 16th International Conference on Educational Data Mining."},{"issue":"12","key":"10.1007\/s40593-025-00494-6_bib217","doi-asserted-by":"crossref","first-page":"2306","DOI":"10.3923\/itj.2013.2306.2314","article-title":"Reinforcement learning approach for adaptive e-learning systems using learning styles","volume":"12","author":"Velusamy","year":"2013","journal-title":"Information Technology Journal"},{"key":"10.1007\/s40593-025-00494-6_bib218","series-title":"2018 International CET Conference on Control, Communication, and Computing (IC4)","first-page":"361","article-title":"A framework for intelligent learning assistant platform based on cognitive computing for children with autism spectrum disorder","author":"Vijayan","year":"2018"},{"key":"10.1007\/s40593-025-00494-6_bib219","series-title":"2023 IEEE International Conference on Advanced Learning Technologies (ICALT)","first-page":"55","article-title":"Learning path recommendation based on knowledge tracing and reinforcement learning","author":"Wan","year":"2023"},{"key":"10.1007\/s40593-025-00494-6_bib220","series-title":"Proceedings of the 6th International Conference on Computer Supported Education, Barcelona, Spain, 1 - 3 April, 2014","first-page":"233","article-title":"Pomdp framework for building an intelligent tutoring system","author":"Wang","year":"2014"},{"key":"10.1007\/s40593-025-00494-6_bib221","series-title":"Lecture Notes in Electrical Engineering: Vol. 276. Future information technology: Futuretech 2013","first-page":"191","article-title":"Learning teaching in teaching: Online reinforcement learning for intelligent tutoring","author":"Wang","year":"2014"},{"issue":"1","key":"10.1007\/s40593-025-00494-6_bib222","doi-asserted-by":"crossref","DOI":"10.1609\/aaai.v32i1.11981","article-title":"MathDQN: Solving arithmetic word problems via deep reinforcement learning","volume":"32","author":"Wang","year":"2018","journal-title":"Proceedings of the AAAI Conference on Artificial Intelligence"},{"issue":"8","key":"10.1007\/s40593-025-00494-6_bib223","doi-asserted-by":"crossref","first-page":"553","DOI":"10.18178\/ijiet.2018.8.8.1098","article-title":"Reinforcement learning in a pomdp based intelligent tutoring system for optimizing teaching strategies","volume":"8","author":"Wang","year":"2018","journal-title":"International Journal of Information and Education Technology"},{"key":"10.1007\/s40593-025-00494-6_bib224","series-title":"Proceedings of 2020 IEEE International Conference on Teaching, Assessment, and Learning for Engineering (TALE): Date and venue: 8\u201311 December 2020, online","first-page":"474","article-title":"Poem: A personalized online education scheme based on reinforcement learning","author":"Wang","year":"2020"},{"key":"10.1007\/s40593-025-00494-6_bib225","series-title":"2023 IEEE International Conference on Bioinformatics and Biomedicine (BIBM)","first-page":"3480","article-title":"Learning path design on knowledge graph by using reinforcement learning","author":"Wang","year":"2023"},{"key":"10.1007\/s40593-025-00494-6_bib226","series-title":"2024 Cross Strait Radio Science and Wireless Technology Conference (CSRSWTC)","first-page":"1","article-title":"Real-Time Wireless Adaptive Learning Systems Using Reinforcement Learning and IoT for Smart Education","author":"Wang","year":"2024"},{"issue":"3\u20134","key":"10.1007\/s40593-025-00494-6_bib227","first-page":"279","article-title":"Q-learning","volume":"8","author":"Watkins","year":"1992","journal-title":"Machine Learning"},{"issue":"2","key":"10.1007\/s40593-025-00494-6_bib228","doi-asserted-by":"crossref","first-page":"152","DOI":"10.1109\/TLT.2017.2692761","article-title":"Approximately optimal teaching of approximately optimal learners","volume":"11","author":"Whitehill","year":"2018","journal-title":"IEEE Transactions on Learning Technologies"},{"issue":"3\u20134","key":"10.1007\/s40593-025-00494-6_bib229","doi-asserted-by":"crossref","first-page":"229","DOI":"10.1023\/A:1022672621406","article-title":"Simple statistical gradient-following algorithms for connectionist reinforcement learning","volume":"8","author":"Williams","year":"1992","journal-title":"Machine Learning"},{"key":"10.1007\/s40593-025-00494-6_bib230","doi-asserted-by":"crossref","first-page":"691","DOI":"10.1109\/TLT.2023.3326449","article-title":"Contrastive Personalized Exercise Recommendation With Reinforcement Learning","volume":"17","author":"Wu","year":"2024","journal-title":"IEEE Transactions on Learning Technologies"},{"key":"10.1007\/s40593-025-00494-6_bib231","series-title":"2024 International Conference on Integrated Intelligence and Communication Systems (ICIICS)","first-page":"1","article-title":"English Learning Knowledge Point Recommendation Algorithm based on Deep Deterministic Policy Gradient","author":"Yang","year":"2024"},{"key":"10.1007\/s40593-025-00494-6_bib232","series-title":"2024 Cross Strait Radio Science and Wireless Technology Conference (CSRSWTC)","first-page":"1","article-title":"Enhancing Student Engagement in Smart Classrooms Using Reinforcement Learning Algorithms","author":"Yantao","year":"2024"},{"key":"10.1007\/s40593-025-00494-6_bib233","series-title":"Lecture Notes in Computer Science: Vol. 13891. Augmented Intelligence and Intelligent Tutoring Systems: 19th International Conference, ITS 2023 Proceedings (1st ed. 2023, Vol. 13891","first-page":"247","article-title":"Using the ITS components in improving the q-learning policy for instructional sequencing","author":"Yessad","year":"2023"},{"key":"10.1007\/s40593-025-00494-6_bib234","series-title":"2024 American Control Conference (ACC)","first-page":"3223","article-title":"Using Reward Shaping to Train Cognitive-Based Control Policies for Intelligent Tutoring Systems","author":"Yuh","year":"2024"},{"key":"10.1007\/s40593-025-00494-6_bib235","series-title":"NeurIPS 2023: IMOL Workshop \"Intrinsically-Motivated and Open-Ended Learning\u201d","first-page":"423","article-title":"Emergence of a Symbolic Goal Representation with an Intelligent Tutoring System based on Intrinsic Motivation","author":"Zadem","year":"2023"},{"key":"10.1007\/s40593-025-00494-6_bib236","series-title":"AAMAS \u201919, Proceedings of the 18th International Conference on Autonomous Agents and MultiAgent Systems","first-page":"711","article-title":"Bootstrapped policy gradient for difficulty adaptation in intelligent tutoring systems","author":"Zhang","year":"2019"},{"key":"10.1007\/s40593-025-00494-6_bib237","series-title":"Springer eBook Collection. Deep Reinforcement Learning: Fundamentals, Research and Applications","first-page":"125","article-title":"Taxonomy of reinforcement learning algorithms","author":"Zhang","year":"2020"},{"key":"10.1007\/s40593-025-00494-6_bib238","series-title":"2023 International Conference on Intelligent Computing, Communication & Convergence (ICI3C)","first-page":"349","article-title":"Game Design and Learning Effectiveness Evaluation of English Teaching Based on Reinforcement Learning Algorithm","author":"Zhang","year":"2023"},{"key":"10.1007\/s40593-025-00494-6_bib239","series-title":"Proceedings of the International Conference on Decision Science & Management","first-page":"179","article-title":"Using deep Reinforcement Learning to Optimize the Motivational Incentive Mechanism of Online English Learners","author":"Zhang","year":"2024"},{"issue":"4","key":"10.1007\/s40593-025-00494-6_bib240","doi-asserted-by":"crossref","first-page":"753","DOI":"10.1007\/s11257-021-09292-w","article-title":"Personalized task difficulty adaptation based on reinforcement learning","volume":"31","author":"Zhang","year":"2021","journal-title":"User Modeling and User-Adapted Interaction"},{"key":"10.1007\/s40593-025-00494-6_bib241","series-title":"Proceedings of the 5th International Conference on I-SMAC (IoT in Social, Mobile, Analytics and Cloud): I-SMAC 2021 : 11\u201313, November 2021","first-page":"414","article-title":"Allocation of english remote guiding based on deep reinforcement learning and multi-objective optimization","author":"Zhiyong","year":"2021"},{"key":"10.1007\/s40593-025-00494-6_bib242","series-title":"LNCS sublibrary: 11625\u201311626. Artificial intelligence in education: 20th international conference, AIED 2019, Chicago, IL, USA, June 25\u201329, 2019, proceedings","first-page":"544","article-title":"Hierarchical reinforcement learning for pedagogical policy induction","author":"Zhou","year":"2019"},{"key":"10.1007\/s40593-025-00494-6_bib243","series-title":"ACM Digital Library, Proceedings of the 28th ACM Conference on User Modeling, Adaptation and Personalization","first-page":"284","article-title":"Improving student-system interaction through data-driven explanations of hierarchical reinforcement learning induced pedagogical policies","author":"Zhou","year":"2020"},{"issue":"2","key":"10.1007\/s40593-025-00494-6_bib244","doi-asserted-by":"crossref","first-page":"454","DOI":"10.1007\/s40593-021-00269-9","article-title":"Leveraging granularity: Hierarchical reinforcement learning for pedagogical policy induction","volume":"32","author":"Zhou","year":"2021","journal-title":"International Journal of Artificial Intelligence in Education"}],"container-title":["International Journal of Artificial Intelligence in Education"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s40593-025-00494-6.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s40593-025-00494-6","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/api.elsevier.com\/content\/article\/PII:S1560429226000612?httpAccept=text\/xml","content-type":"text\/xml","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/api.elsevier.com\/content\/article\/PII:S1560429226000612?httpAccept=text\/plain","content-type":"text\/plain","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s40593-025-00494-6.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,5,18]],"date-time":"2026-05-18T06:54:23Z","timestamp":1779087263000},"score":1,"resource":{"primary":{"URL":"https:\/\/linkinghub.elsevier.com\/retrieve\/pii\/S1560429226000612"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,12]]},"references-count":244,"journal-issue":{"issue":"5","published-print":{"date-parts":[[2025,12]]}},"alternative-id":["S1560429226000612"],"URL":"https:\/\/doi.org\/10.1007\/s40593-025-00494-6","relation":{},"ISSN":["1560-4292"],"issn-type":[{"value":"1560-4292","type":"print"}],"subject":[],"published":{"date-parts":[[2025,12]]},"assertion":[{"value":"Elsevier","name":"publisher","label":"This article is maintained by"},{"value":"Reinforcement Learning in Education: A Systematic Literature Review","name":"articletitle","label":"Article Title"},{"value":"International Journal of Artificial Intelligence in Education","name":"journaltitle","label":"Journal Title"},{"value":"https:\/\/doi.org\/10.1007\/s40593-025-00494-6","name":"articlelink","label":"CrossRef DOI link to publisher maintained version"},{"value":"article","name":"content_type","label":"Content Type"},{"value":"Copyright \u00a9 2025 The Author(s). Published by Elsevier Ltd","name":"copyright","label":"Copyright"}]}}