{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,6,24]],"date-time":"2026-06-24T17:11:30Z","timestamp":1782321090191,"version":"3.54.5"},"reference-count":43,"publisher":"Springer Science and Business Media LLC","issue":"11","license":[{"start":{"date-parts":[[2025,10,15]],"date-time":"2025-10-15T00:00:00Z","timestamp":1760486400000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/creativecommons.org\/licenses\/by\/4.0"},{"start":{"date-parts":[[2025,10,15]],"date-time":"2025-10-15T00:00:00Z","timestamp":1760486400000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/creativecommons.org\/licenses\/by\/4.0"}],"funder":[{"DOI":"10.13039\/501100004462","name":"Consiglio Nazionale Delle Ricerche","doi-asserted-by":"crossref","id":[{"id":"10.13039\/501100004462","id-type":"DOI","asserted-by":"crossref"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Mach Learn"],"published-print":{"date-parts":[[2025,11]]},"abstract":"<jats:title>Abstract<\/jats:title>\n                  <jats:p>This paper presents a methodology for integrating human expert knowledge into machine learning (ML) workflows to improve both model interpretability and the quality of explanations produced by explainable AI (XAI) techniques. We strive to enhance standard ML and XAI pipelines without modifying underlying algorithms, focusing instead on embedding domain knowledge at two stages: (1) during model development through expert-guided data structuring and feature engineering, and (2) during explanation generation via domain-aware synthetic neighbourhoods. Visual analytics is used to support experts in transforming raw data into semantically richer representations. We validate the methodology in two case studies: predicting COVID-19 incidence and classifying vessel movement patterns. The studies demonstrated improved alignment of models with expert reasoning and better quality of synthetic neighbourhoods. We also explore using large language models (LLMs) to assist experts in developing domain-compliant data generators. Our findings highlight both the benefits and limitations of existing XAI methods and point to a research direction for addressing these gaps.<\/jats:p>","DOI":"10.1007\/s10994-025-06879-x","type":"journal-article","created":{"date-parts":[[2025,10,15]],"date-time":"2025-10-15T19:59:46Z","timestamp":1760558386000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":3,"title":["Integrating human knowledge for explainable AI"],"prefix":"10.1007","volume":"114","author":[{"given":"Eleonora","family":"Cappuccio","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Bahavathy","family":"Kathirgamanathan","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Salvatore","family":"Rinzivillo","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Gennady","family":"Andrienko","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Natalia","family":"Andrienko","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2025,10,15]]},"reference":[{"key":"6879_CR1","doi-asserted-by":"publisher","unstructured":"Abdul, A. M., Vermeulen, J., Wang, D., Lim, B. Y., & Kankanhalli, M. S. (2018). Trends and trajectories for explainable, accountable and intelligible systems: An HCI research agenda. ACM. https:\/\/doi.org\/10.1145\/3173574.3174156","DOI":"10.1145\/3173574.3174156"},{"key":"6879_CR2","doi-asserted-by":"publisher","unstructured":"Adadi, A., & Berrada, M. (2018). Peeking inside the black-box: A survey on explainable artificial intelligence (XAI). IEEE Access. https:\/\/doi.org\/10.1109\/ACCESS.2018.2870052","DOI":"10.1109\/ACCESS.2018.2870052"},{"key":"6879_CR3","doi-asserted-by":"publisher","unstructured":"Andrienko, N., Andrienko, G., Artikis, A., Mantenoglou, P., & Rinzivillo, S. (2024). Human-in-the-loop: Visual analytics for building models recognising behavioural patterns in time series. IEEE Computer Graphics and Applications. https:\/\/doi.org\/10.1109\/MCG.2024.3379851","DOI":"10.1109\/MCG.2024.3379851"},{"key":"6879_CR4","doi-asserted-by":"publisher","first-page":"172","DOI":"10.1016\/j.is.2015.08.007","volume":"57","author":"N Andrienko","year":"2016","unstructured":"Andrienko, N., Andrienko, G., & Rinzivillo, S. (2016). Leveraging spatial abstraction in traffic analysis and forecasting with visual analytics. Information Systems, 57, 172\u2013194. https:\/\/doi.org\/10.1016\/j.is.2015.08.007","journal-title":"Information Systems"},{"key":"6879_CR5","doi-asserted-by":"publisher","unstructured":"Andrienko, N., Andrienko, G., & Shirato, G. (2023). Episodes and topics in multivariate temporal data,42(6), Article e14926. https:\/\/doi.org\/10.1111\/cgf.14926","DOI":"10.1111\/cgf.14926"},{"key":"6879_CR6","doi-asserted-by":"publisher","unstructured":"Bhattacharya, A., Verbert, K., et\u00a0al. (2024). Exmos: Explanatory model steering through multifaceted explanations and data configurations. In: Proceedings of the CHI Conference on Human Factors in Computing Systems. ACM. https:\/\/doi.org\/10.1145\/3544548.3581135","DOI":"10.1145\/3544548.3581135"},{"key":"6879_CR7","doi-asserted-by":"crossref","unstructured":"Buchanan, B.G., Davis, R., & Feigenbaum, E.A. (2006). Expert systems: A perspective from computer science. The Cambridge handbook of expertise and expert performance pp. 87\u2013103.","DOI":"10.1017\/CBO9780511816796.006"},{"key":"6879_CR8","doi-asserted-by":"publisher","unstructured":"Cheng, F., Ming, Y., & Qu, H. (2021). DECE: decision explorer with counterfactual explanations for machine learning models. IEEE. https:\/\/doi.org\/10.1109\/TVCG.2020.3030342","DOI":"10.1109\/TVCG.2020.3030342"},{"key":"6879_CR9","doi-asserted-by":"publisher","unstructured":"Dong, G., & Liu, H. (2018). Feature engineering for machine learning and data analytics. CRC Press. https:\/\/doi.org\/10.1201\/9781315181080","DOI":"10.1201\/9781315181080"},{"key":"6879_CR10","doi-asserted-by":"publisher","unstructured":"Endert, A., Ribarsky, W., Turkay, C., Wong, B. W., Nabney, I., Blanco, I. D., & Rossi, F. (2017). The state of the art in integrating machine learning into visual analytics. Computer Graphics Forum. https:\/\/doi.org\/10.1111\/cgf.13092","DOI":"10.1111\/cgf.13092"},{"key":"6879_CR11","doi-asserted-by":"publisher","unstructured":"Giabbanelli, P. J., & Jackson, P. J. (2015). Using visual analytics to support the integration of expert knowledge in the design of medical models and simulations. Procedia Computer Science. https:\/\/doi.org\/10.1016\/j.procs.2015.05.195","DOI":"10.1016\/j.procs.2015.05.195"},{"issue":"25","key":"6879_CR12","first-page":"723","volume":"13","author":"A Gretton","year":"2012","unstructured":"Gretton, A., Borgwardt, K. M., Rasch, M. J., Sch\u00f6lkopf, B., & Smola, A. (2012). A kernel two-sample test. Journal of Machine Learning Research, 13(25), 723\u2013773.","journal-title":"Journal of Machine Learning Research"},{"key":"6879_CR13","doi-asserted-by":"publisher","unstructured":"Guidotti, R., Monreale, A., Giannotti, F., Pedreschi, D., Ruggieri, S., & Turini, F. (2019). Factual and counterfactual explanations for black box decision making. IEEE Intelligent System https:\/\/doi.org\/10.1109\/MIS.2019.2957223","DOI":"10.1109\/MIS.2019.2957223"},{"key":"6879_CR14","doi-asserted-by":"publisher","unstructured":"Guidotti, R., Monreale, A., Ruggieri, S., Turini, F., Giannotti, F., & Pedreschi, D. (2019). A survey of methods for explaining black box models. ACM Computer Survery https:\/\/doi.org\/10.1145\/3236009","DOI":"10.1145\/3236009"},{"key":"6879_CR15","doi-asserted-by":"publisher","unstructured":"Hohman, F., Kahng, M., Pienta, R. S., & Chau, D. H. (2019). Visual analytics in deep learning: An interrogative survey for the next frontiers. IEEE Transactions on Visualization and Computer Graphics https:\/\/doi.org\/10.1109\/TVCG.2018.2843369","DOI":"10.1109\/TVCG.2018.2843369"},{"key":"6879_CR16","doi-asserted-by":"publisher","unstructured":"Karpatne, A., Atluri, G., Faghmous, J. H., Steinbach, M. S., Banerjee, A., Ganguly, A. R., Shekhar, S., Samatova, N. F., & Kumar, V. (2017). Theory-guided data science: A new paradigm for scientific discovery from data. IEEE Transactions on Knowledge and Data Engineering https:\/\/doi.org\/10.1109\/TKDE.2017.2720168","DOI":"10.1109\/TKDE.2017.2720168"},{"key":"6879_CR17","doi-asserted-by":"publisher","unstructured":"Kidd, A.L. (ed.). (1987). Knowledge Acquisition for Expert Systems: A Practical Handbook. Springer, New York, NY. https:\/\/doi.org\/10.1007\/978-1-4613-1823-1","DOI":"10.1007\/978-1-4613-1823-1"},{"key":"6879_CR18","doi-asserted-by":"publisher","unstructured":"Kulesza, T., Stumpf, S., Burnett, M. M., Yang, S., Kwan, I., & Wong, W. (2013). Too much, too little, or just right? ways explanations impact end users\u2019 mental models. IEEE Computer Society. https:\/\/doi.org\/10.1109\/VLHCC.2013.6645235","DOI":"10.1109\/VLHCC.2013.6645235"},{"key":"6879_CR19","doi-asserted-by":"publisher","unstructured":"Ponce-de Leon, M., del Valle, J., Fernandez, J. M., Bernardo, M., Cirillo, D., Sanchez-Valle, J., Smith, M., Capella-Gutierrez, S., Gull\u00f3n, T., & Valencia, A. (2021). COVID-19 Flow-Maps an open geographic information system on COVID-19 and human mobility for Spain. Scientific Data. https:\/\/doi.org\/10.1038\/s41597-021-01093-5","DOI":"10.1038\/s41597-021-01093-5"},{"key":"6879_CR20","doi-asserted-by":"publisher","unstructured":"Ponce-de Leon, M., del Valle, J., Fern\u00e1ndez, J.M., Bernardo, M., Cirillo, D., Sanchez-Valle, J., Smith, M., Capella-Gutierrez, S., Gull\u00f3n, T., & Valencia, A. (2021). COVID19 Flow-Maps daily cases reports.https:\/\/doi.org\/10.5281\/zenodo.5217386","DOI":"10.5281\/zenodo.5217386"},{"key":"6879_CR21","doi-asserted-by":"publisher","unstructured":"Ponce-de Leon, M., del Valle, J., Fern\u00e1ndez, J.M., Bernardo, M., Cirillo, D., Sanchez-Valle, J., Smith, M., Capella-Gutierrez, S., Gull\u00f3n, T., & Valencia, A. (2021). COVID19 Flow-Maps daily-mobility for Spain.https:\/\/doi.org\/10.5281\/zenodo.5539411","DOI":"10.5281\/zenodo.5539411"},{"key":"6879_CR22","doi-asserted-by":"publisher","unstructured":"Ponce-de Leon, M., del Valle, J., Fern\u00e1ndez, J.M., Bernardo, M., Cirillo, D., Sanchez-Valle, J., Smith, M., Capella-Gutierrez, S., Gull\u00f3n, T., & Valencia, A. (2021). COVID19 Flow-Maps population data. https:\/\/doi.org\/10.5281\/zenodo.5226351","DOI":"10.5281\/zenodo.5226351"},{"key":"6879_CR23","unstructured":"Liao, Q.V., & Varshney, K.R. (2021). Human-centered explainable AI (XAI): from algorithms to user experiences. CoRR arxiv:2110.10790"},{"key":"6879_CR24","doi-asserted-by":"publisher","unstructured":"Lu, Y., Garcia, R., Hansen, B., Gleicher, M., & Maciejewski, R. (2017). The state-of-the-art in predictive visual analytics. Computer Graphics Forum. https:\/\/doi.org\/10.1111\/cgf.13210","DOI":"10.1111\/cgf.13210"},{"key":"6879_CR25","volume-title":"A unified approach to interpreting model predictions.","author":"SM Lundberg","year":"2017","unstructured":"Lundberg, S. M., & Lee, S. I. (2017). A unified approach to interpreting model predictions. Red Hook, NY, USA: NIPS\u201917, Curran Associates Inc."},{"key":"6879_CR26","doi-asserted-by":"publisher","unstructured":"Miller, T. (2019). Explanation in artificial intelligence: Insights from the social sciences. Artificail Intelligence https:\/\/doi.org\/10.1016\/J.ARTINT.2018.07.007","DOI":"10.1016\/J.ARTINT.2018.07.007"},{"key":"6879_CR27","doi-asserted-by":"publisher","unstructured":"Ming, Y., Qu, H., & Bertini, E. (2019). Rulematrix: Visualizing and understanding classifiers with rules. IEEE Transactions on Visualization and Computer Graphics, https:\/\/doi.org\/10.1109\/TVCG.2018.2864812","DOI":"10.1109\/TVCG.2018.2864812"},{"key":"6879_CR28","doi-asserted-by":"publisher","unstructured":"Muralidhar, N., Islam, M.R., Marwah, M., Karpatne, A., & Ramakrishnan, N. (2018). Incorporating prior domain knowledge into deep neural networks. In: 2018 IEEE International Conference on Big Data (Big Data). pp. 36\u201345. https:\/\/doi.org\/10.1109\/BigData.2018.8621955","DOI":"10.1109\/BigData.2018.8621955"},{"issue":"2","key":"6879_CR29","doi-asserted-by":"publisher","DOI":"10.1016\/j.visinf.2025.100234","volume":"9","author":"L Peng","year":"2025","unstructured":"Peng, L., Lin, Z., Andrienko, N., Andrienko, G., & Chen, S. (2025). Contextualized visual analytics for multivariate events. Visual Informatics, 9(2), Article 100234. https:\/\/doi.org\/10.1016\/j.visinf.2025.100234","journal-title":"Visual Informatics"},{"key":"6879_CR30","doi-asserted-by":"publisher","unstructured":"Raissi, M., Perdikaris, P., & Karniadakis, G. (2019). Physics-informed neural networks: A deep learning framework for solving forward and inverse problems involving nonlinear partial differential equations. Journal of Computational Physics. https:\/\/doi.org\/10.1016\/j.jcp.2018.10.045","DOI":"10.1016\/j.jcp.2018.10.045"},{"key":"6879_CR31","doi-asserted-by":"publisher","unstructured":"Rajabi, E., & Etminani, K. (2024). Knowledge-graph-based explainable AI: A systematic review.  Journal of Information Science https:\/\/doi.org\/10.1177\/01655515221112844","DOI":"10.1177\/01655515221112844"},{"key":"6879_CR32","doi-asserted-by":"publisher","unstructured":"Ribeiro, M. T., Singh, S., & Guestrin, C. (2016). \u201cWhy Should I Trust You?\u201d: Explaining the Predictions of Any Classifier. ACM. https:\/\/doi.org\/10.1145\/2939672.2939778","DOI":"10.1145\/2939672.2939778"},{"key":"6879_CR33","doi-asserted-by":"publisher","unstructured":"Ribeiro, M.T., Singh, S., & Guestrin, C. (2018). Anchors: High-precision model-agnostic explanations. In: McIlraith, S.A., Weinberger, K.Q. (eds.) Proc. of the Thirty-Second AAAI Conference on Artificial Intelligence. AAAI Press. https:\/\/doi.org\/10.1609\/AAAI.V32I1.11491","DOI":"10.1609\/AAAI.V32I1.11491"},{"key":"6879_CR34","doi-asserted-by":"publisher","unstructured":"Riveiro, M., & Thill, S. (2021). \u201cthat\u2019s (not) the output I expected!\u201d on the role of end user expectations in creating explanations of AI systems. Artificail Intelligence. https:\/\/doi.org\/10.1016\/J.ARTINT.2021.103507","DOI":"10.1016\/J.ARTINT.2021.103507"},{"key":"6879_CR35","doi-asserted-by":"publisher","unstructured":"von Rueden, L., Mayer, S., Beckh, K., Georgiev, B., Giesselbach, S., Heese, R., Kirsch, B., Pfrommer, J., Pick, A., Ramamurthy, R., Walczak, M., Garcke, J., Bauckhage, C., & Schuecker, J. (2023). Informed machine learning \u2013 a taxonomy and survey of integrating prior knowledge into learning systems. IEEE Transactions on Knowledge and Data Engineering. https:\/\/doi.org\/10.1109\/TKDE.2021.3079836","DOI":"10.1109\/TKDE.2021.3079836"},{"key":"6879_CR36","doi-asserted-by":"publisher","unstructured":"Sacha, D., Kraus, M., Keim, D. A., & Chen, M. (2019). Vis4ml: An ontology for visual analytics assisted machine learning. IEEE Transactions on Visualization and Computer Graphics. https:\/\/doi.org\/10.1109\/TVCG.2018.2864838","DOI":"10.1109\/TVCG.2018.2864838"},{"issue":"8","key":"6879_CR37","doi-asserted-by":"publisher","first-page":"476","DOI":"10.1038\/s42256-020-0212-3","volume":"2","author":"P Schramowski","year":"2020","unstructured":"Schramowski, P., Stammer, W., Teso, S., Brugger, A., Herbert, F., Shao, X., Luigs, L., & Kersting, K. (2020). Making deep neural networks right for the right scientific reasons by interacting with their explanations. Nature Machine Intelligence, 2(8), 476\u2013486. https:\/\/doi.org\/10.1038\/s42256-020-0212-3","journal-title":"Nature Machine Intelligence"},{"key":"6879_CR38","doi-asserted-by":"publisher","unstructured":"Simkute, A., Surana, A., Luger, E., Evans, M., & Jones, R. (2022). XAI for learning: Narrowing down the digital divide between \u201cnew\u201d and \u201cold\u201d experts. ACM. https:\/\/doi.org\/10.1145\/3547522.3547678","DOI":"10.1145\/3547522.3547678"},{"key":"6879_CR39","doi-asserted-by":"publisher","unstructured":"Spinner, T., Schlegel, U., Sch\u00e4fer, H., & El-Assady, M. (2020). explAIner: A visual analytics framework for interactive and explainable machine learning. IEEE Transactions on Visualization and Computer Graphics. https:\/\/doi.org\/10.1109\/TVCG.2019.2934629","DOI":"10.1109\/TVCG.2019.2934629"},{"issue":"8","key":"6879_CR40","doi-asserted-by":"publisher","first-page":"1249","DOI":"10.1016\/j.jspi.2013.03.018","volume":"143","author":"GJ Sz\u00e9kely","year":"2013","unstructured":"Sz\u00e9kely, G. J., & Rizzo, M. L. (2013). Energy statistics: A class of statistics based on distances. Journal of Statistical Planning and Inference, 143(8), 1249\u20131272. https:\/\/doi.org\/10.1016\/j.jspi.2013.03.018https:\/\/www.sciencedirect.com\/science\/article\/pii\/S0378375813000633.","journal-title":"Journal of Statistical Planning and Inference"},{"key":"6879_CR41","doi-asserted-by":"publisher","unstructured":"Teso, S., & Kersting, K. (2019). Explanatory interactive machine learning. In: Proceedings of the 2019 AAAI\/ACM Conference on AI, Ethics, and Society. pp. 239\u2013245. ACM. https:\/\/doi.org\/10.1145\/3306618.3314293","DOI":"10.1145\/3306618.3314293"},{"key":"6879_CR42","doi-asserted-by":"publisher","unstructured":"Wang, J., Liu, S., & Zhang, W. (2024). Visual analytics for machine learning: A data perspective survey. IEEE Transactions on Visualization and Computer Graphics. https:\/\/doi.org\/10.1109\/TVCG.2024.3357065","DOI":"10.1109\/TVCG.2024.3357065"},{"key":"6879_CR43","doi-asserted-by":"publisher","unstructured":"Zhao, J., Glueck, M., Isenberg, P., Chevalier, F., & Khan, A. (2018). Supporting handoff in asynchronous collaborative sensemaking using knowledge-transfer graphs. IEEE Transactions on Visualization and Computer Graphics. https:\/\/doi.org\/10.1109\/TVCG.2017.2745279","DOI":"10.1109\/TVCG.2017.2745279"}],"container-title":["Machine Learning"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10994-025-06879-x.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s10994-025-06879-x\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10994-025-06879-x.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,12,2]],"date-time":"2025-12-02T14:29:06Z","timestamp":1764685746000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s10994-025-06879-x"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,10,15]]},"references-count":43,"journal-issue":{"issue":"11","published-print":{"date-parts":[[2025,11]]}},"alternative-id":["6879"],"URL":"https:\/\/doi.org\/10.1007\/s10994-025-06879-x","relation":{},"ISSN":["0885-6125","1573-0565"],"issn-type":[{"value":"0885-6125","type":"print"},{"value":"1573-0565","type":"electronic"}],"subject":[],"published":{"date-parts":[[2025,10,15]]},"assertion":[{"value":"3 April 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"7 July 2025","order":2,"name":"revised","label":"Revised","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"25 August 2025","order":3,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"15 October 2025","order":4,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare no competing interests.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}],"article-number":"250"}}