{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,11]],"date-time":"2026-07-11T11:12:02Z","timestamp":1783768322758,"version":"3.55.0"},"reference-count":42,"publisher":"Springer Science and Business Media LLC","issue":"4","license":[{"start":{"date-parts":[[2026,5,8]],"date-time":"2026-05-08T00:00:00Z","timestamp":1778198400000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,5,8]],"date-time":"2026-05-08T00:00:00Z","timestamp":1778198400000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Multimedia Systems"],"published-print":{"date-parts":[[2026,6]]},"DOI":"10.1007\/s00530-026-02390-y","type":"journal-article","created":{"date-parts":[[2026,5,8]],"date-time":"2026-05-08T01:44:45Z","timestamp":1778204685000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["Multimodal fake news video contradiction explanations with large language models"],"prefix":"10.1007","volume":"32","author":[{"given":"Kuangda","family":"Hu","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Yan","family":"Liao","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Hao","family":"Liu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Hengxuan","family":"Lin","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Zikun","family":"Hu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2026,5,8]]},"reference":[{"key":"2390_CR1","doi-asserted-by":"crossref","unstructured":"Wittenberg, C., Tappin, B.M., Berinsky, A.J., et al.: The (minimal) persuasive advantage of political video over text. Proc. Natl. Acad. Sci. 118(47), e2114388118 (2021)","DOI":"10.1073\/pnas.2114388118"},{"key":"2390_CR2","unstructured":"ElBoghdady, D.: Market quavers after fake AP tweet says Obama was hurt in White House explosions. Wash. Post 23 (2013)"},{"issue":"3","key":"2390_CR3","doi-asserted-by":"publisher","first-page":"233","DOI":"10.1111\/hir.12311","volume":"37","author":"SB Naeem","year":"2020","unstructured":"Naeem, S.B., Bhatti, R.: The Covid-19 \u2018infodemic\u2019: a new front for information professionals. Health Inform. Libr. J. 37(3), 233\u2013239 (2020)","journal-title":"Health Inform. Libr. J."},{"key":"2390_CR4","doi-asserted-by":"crossref","unstructured":"Qian, S., Wang, J., Hu, J., et al.: Hierarchical multi-modal contextual attention network for fake news detection[C]\/\/Proceedings of the 44th international ACM SIGIR conference on research and development in information retrieval. 153\u2013162 (2021)","DOI":"10.1145\/3404835.3462871"},{"key":"2390_CR5","doi-asserted-by":"crossref","unstructured":"Zhou, Y., Yang, Y., Ying, Q., et al.: Multi-modal fake news detection on social media via multi-grained information fusion[C]\/\/Proceedings of the 2023 ACM International Conference on Multimedia Retrieval 343\u2013352. (2023)","DOI":"10.1145\/3591106.3592271"},{"issue":"6","key":"2390_CR6","doi-asserted-by":"publisher","first-page":"301","DOI":"10.1093\/jcmc\/zmab010","volume":"26","author":"SS Sundar","year":"2021","unstructured":"Sundar, S.S., Molina, M.D., Cho, E.: Seeing is believing: Is video modality more powerful in spreading fake news via online messaging apps? J. Comput. Mediat. Commun. 26(6), 301\u2013319 (2021)","journal-title":"J. Comput. Mediat. Commun."},{"key":"2390_CR7","doi-asserted-by":"crossref","unstructured":"Hou, R., P\u00e9rez-Rosas, V., Loeb, S., et al.: Towards automatic detection of misinformation in online medical videos[C]\/\/2019 International conference on multimodal interaction 235\u2013243 (2019)","DOI":"10.1145\/3340555.3353763"},{"key":"2390_CR8","doi-asserted-by":"crossref","unstructured":"Choi, H., Ko, Y.: Using topic modeling and adversarial neural networks for fake news video detection[C]\/\/Proceedings of the 30th ACM international conference on information & knowledge management 2950\u20132954 (2021)","DOI":"10.1145\/3459637.3482212"},{"key":"2390_CR9","doi-asserted-by":"crossref","unstructured":"Shang, L., Kou, Z., Zhang, Y., et al.: A multimodal misinformation detector for covid-19 short videos on tiktok[C]\/\/2021 IEEE international conference on big data (big data). IEEE 899\u2013908 (2021)","DOI":"10.1109\/BigData52589.2021.9671928"},{"key":"2390_CR10","doi-asserted-by":"crossref","unstructured":"Bu, Y., Sheng, Q., Cao, J., et al.: Combating online misinformation videos: characterization, detection, and future directions[C]\/\/Proceedings of the 31st ACM International Conference on Multimedia 8770\u20138780 (2023)","DOI":"10.1145\/3581783.3612426"},{"key":"2390_CR11","unstructured":"Venkatagiri, S., Schafer, J.S., Prochaska, S.: The Challenges of Studying Misinformation on Video-Sharing Platforms During Crises and Mass-Convergence Events. arxiv preprint arxiv:2303.14309 (2023)"},{"key":"2390_CR12","doi-asserted-by":"crossref","unstructured":"Byrne, R.M.J.: Counterfactuals in explainable artificial intelligence (XAI): Evidence from human reasoning[C]\/\/IJCAI. 6276\u20136282 (2019)","DOI":"10.24963\/ijcai.2019\/876"},{"key":"2390_CR13","doi-asserted-by":"crossref","unstructured":"Keane, M.T., Kenny, E.M., Delaney, E., et al.: If only we had better counterfactual explanations: five key deficits to rectify in the evaluation of counterfactual xai techniques. arXiv preprint arXiv:2103.01035 (2021)","DOI":"10.24963\/ijcai.2021\/609"},{"key":"2390_CR14","doi-asserted-by":"publisher","first-page":"22199","DOI":"10.52202\/068431-1613","volume":"35","author":"T Kojima","year":"2022","unstructured":"Kojima, T., Gu, S.S., Reid, M., et al.: Large language models are zero-shot reasoners. Adv. Neural. Inf. Process. Syst. 35, 22199\u201322213 (2022)","journal-title":"Adv. Neural. Inf. Process. Syst."},{"key":"2390_CR15","unstructured":"Luo, L., Zhang, G., Xu, H., et al.: End-to-end neuro-symbolic reinforcement learning with textual explanations[C]\/\/Forty-first International Conference on Machine Learning"},{"key":"2390_CR16","unstructured":"Chen, Z., Singh, A.K., Sra, M.: LMExplainer: a Knowledge-Enhanced Explainer for Language Models. arXiv preprint arXiv:2303.16537 (2023)"},{"key":"2390_CR17","doi-asserted-by":"publisher","first-page":"24824","DOI":"10.52202\/068431-1800","volume":"35","author":"J Wei","year":"2022","unstructured":"Wei, J., Wang, X., Schuurmans, D., et al.: Chain-of-thought prompting elicits reasoning in large language models. Adv. Neural. Inf. Process. Syst. 35, 24824\u201324837 (2022)","journal-title":"Adv. Neural. Inf. Process. Syst."},{"key":"2390_CR18","unstructured":"Dong, Q., Li, L., Dai, D., et al.: A survey on in-context learning. arXiv preprint arXiv:2301.00234 (2022)"},{"key":"2390_CR19","doi-asserted-by":"crossref","unstructured":"Fuxiao Liu, Y., Yacoob, Shrivastava, A.: COVID-VTS: Fact Extraction and Verification on Short Video Platforms. In Proceedings of the 17th Conference of the European Chapter of the Association for Computational Linguistics 178\u2013188 (2023)","DOI":"10.18653\/v1\/2023.eacl-main.14"},{"key":"2390_CR20","unstructured":"Official-NV: a news video dataset for multimodal fake news detection"},{"key":"2390_CR21","unstructured":"Serrano, J.C.M., Papakyriakopoulos, O., Hegelich, S.: NLP-based feature extraction for the detection of COVID-19 misinformation videos on YouTube. In Proceedings of the 1st Workshop on NLP for COVID-19 at ACL 2020 (2020)"},{"key":"2390_CR22","unstructured":"Jagtap, R., Kumar, A., Goel, R., Sharma, S., Sharma, R., Clint, P.G.: Misinformation detection on youtube using video captions. arXiv:2107.00941 (2021)"},{"key":"2390_CR23","doi-asserted-by":"crossref","unstructured":"Jacob Devlin, M.-W., Chang, K., Lee, Toutanova, K.: BERT: pre-training of deep bidirectional transformers for language understanding. In Proceedings of the 2019 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies, Volume 1 (Long and Short Papers). 4171\u20134186 (2019)","DOI":"10.18653\/v1\/N19-1423"},{"key":"2390_CR24","unstructured":"Wang, K., Chan, D., Zhao, S.Z., Canny, J., and Avideh Zakhor:. Misinformation detection in social media video posts. arXiv:2202.07706 (2022)"},{"key":"2390_CR25","doi-asserted-by":"crossref","unstructured":"Scott McCrae, K., Wang, Zakhor, A.: Multi-modal semantic inconsistency detection in social media news posts. In Multi Media Modeling: MMM 331\u2013343 (2022)","DOI":"10.1007\/978-3-030-98355-0_28"},{"key":"2390_CR26","unstructured":"Alec Radford, J.W., Kim, C., Hallacy, A., Ramesh, G., Goh, S., Agarwal, G., Sastry, A., Askell, P., Mishkin, J., Clark, G., Krueger, and Ilya Sutskever:. Learning transferable visual models from natural language supervision. In Proceedings of the 38th International Conference on Machine Learning 139, 8748\u20138763 (2021)"},{"key":"2390_CR27","unstructured":"Hendricks, L., Anne,14th European, Conference, et al.: Generating visual explanations. Computer Vision\u2013ECCV, Amsterdam, The Netherlands, October 11\u201314, 2016, Proceedings, Part IV 14. Springer International Publishing (2016)"},{"key":"2390_CR28","doi-asserted-by":"crossref","unstructured":"Kim, X., et al.: Textual explanations for self-driving vehicles. Proceedings of the European conference on computer vision (ECCV) (2018)","DOI":"10.1007\/978-3-030-01216-8_35"},{"key":"2390_CR29","doi-asserted-by":"crossref","unstructured":"Kayser, M., Camburu, O.-M., Salewski, L., Emde, C., Do, V., Akata, Z., Lukasiewicz, T.: e-ViL: A dataset and benchmark for natural language explanations in vision-language tasks. arXiv preprint arXiv:2105.03761 (2021)","DOI":"10.1109\/ICCV48922.2021.00128"},{"key":"2390_CR30","doi-asserted-by":"crossref","unstructured":"Byrne, R.M.: J. The rational imagination: How people create alternatives to reality[M]. MIT Press (2007)","DOI":"10.1017\/S0140525X07002579"},{"issue":"3","key":"2390_CR31","doi-asserted-by":"publisher","first-page":"87","DOI":"10.1109\/MIS.2018.033001421","volume":"33","author":"R Hoffman","year":"2018","unstructured":"Hoffman, R., Miller, T., Mueller, S.T., et al.: Explaining explanation, part 4: a deep dive on deep nets. IEEE. Intell. Syst. 33(3), 87\u201395 (2018)","journal-title":"IEEE. Intell. Syst."},{"key":"2390_CR32","doi-asserted-by":"crossref","unstructured":"Mittelstadt, B., Russell, C., Wachter, S.: Explaining explanations in AI[C]\/\/Proceedings of the conference on fairness, accountability, and transparency 279\u2013288 (2019)","DOI":"10.1145\/3287560.3287574"},{"key":"2390_CR33","unstructured":"Li, J., Li, D., Savarese, S., et al.: Blip-2: Bootstrapping language-image pre-training with frozen image encoders and large language models[C]\/\/International conference on machine learning. PMLR 19730\u201319742 (2023)"},{"key":"2390_CR34","doi-asserted-by":"crossref","unstructured":"Zhao, Z., Cao, Y., Gong, S., et al.: Enhancing Zero-Shot Facial Expression Recognition by LLM Knowledge Transfer[J]. (2024). arXiv preprint arXiv:2405.19100","DOI":"10.1109\/WACV61041.2025.00089"},{"key":"2390_CR35","doi-asserted-by":"crossref","unstructured":"Lewis, M., Liu, Y., Goyal, N., et al.: BART: Denoising Sequence-to-Sequence Pre-training for Natural Language Generation, Translation, and Comprehension[C]\/\/Meeting of the Association for Computational Linguistics. Association for Computational Linguistics (2020)","DOI":"10.18653\/v1\/2020.acl-main.703"},{"issue":"70","key":"2390_CR36","first-page":"1","volume":"25","author":"HW Chung","year":"2024","unstructured":"Chung, H.W., Hou, L., Longpre, S., et al.: Scaling instruction-finetuned language models. J. Mach. Learn. Res. 25(70), 1\u201353 (2024)","journal-title":"J. Mach. Learn. Res."},{"key":"2390_CR37","doi-asserted-by":"crossref","unstructured":"Zheng, L., Chiang, W.L., Sheng, Y., et al.: Judging llm-as-a-judge with mt-bench and chatbot arena. Adv. Neural. Inf. Process. Syst. 36 (2024)","DOI":"10.52202\/075280-2020"},{"key":"2390_CR38","doi-asserted-by":"crossref","unstructured":"Reimers, N., Gurevych, I.: Sentence-bert: Sentence embeddings using siamese bert-networks. arXiv preprint arXiv:1908.10084 (2019)","DOI":"10.18653\/v1\/D19-1410"},{"key":"2390_CR39","doi-asserted-by":"crossref","unstructured":"Qi, P., Bu, Y., Cao, J., et al.: Fakesv: A multimodal benchmark with rich social context for fake news detection on short video platforms[C]\/\/Proceedings of the AAAI Conference on Artificial Intelligence 37(12): 14444\u201314452. (2023)","DOI":"10.1609\/aaai.v37i12.26689"},{"key":"2390_CR40","doi-asserted-by":"crossref","unstructured":"Huang, X., Ma, T., Tang, H., et al.: Knowledge-enhanced dynamic scene graph attention network for fake news video detection. IEEE Transactions on Multimedia (2025)","DOI":"10.1109\/TMM.2025.3623491"},{"key":"2390_CR41","doi-asserted-by":"crossref","unstructured":"Shen, J., Wang, Y., Wang, S., et al.: Multi-modal similarity guided adaptive fusion network for short video fake news detection[C]\/\/Proceedings of the 2025 International Conference on Multimedia Retrieval 1145\u20131153 (2025)","DOI":"10.1145\/3731715.3733400"},{"key":"2390_CR42","doi-asserted-by":"crossref","unstructured":"Kayser, M., Camburu, O.M., Salewski, L., et al.: e-vil: A dataset and benchmark for natural language explanations in vision-language tasks[C]\/\/Proceedings of the IEEE\/CVF international conference on computer vision 1244\u20131254 (2021)","DOI":"10.1109\/ICCV48922.2021.00128"}],"container-title":["Multimedia Systems"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00530-026-02390-y.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s00530-026-02390-y","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s00530-026-02390-y.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,7,11]],"date-time":"2026-07-11T10:18:00Z","timestamp":1783765080000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s00530-026-02390-y"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,5,8]]},"references-count":42,"journal-issue":{"issue":"4","published-print":{"date-parts":[[2026,6]]}},"alternative-id":["2390"],"URL":"https:\/\/doi.org\/10.1007\/s00530-026-02390-y","relation":{},"ISSN":["0942-4962","1432-1882"],"issn-type":[{"value":"0942-4962","type":"print"},{"value":"1432-1882","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,5,8]]},"assertion":[{"value":"16 November 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"28 March 2026","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"8 May 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Competing interests"}}],"article-number":"316"}}