{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,1,9]],"date-time":"2026-01-09T02:06:12Z","timestamp":1767924372796,"version":"3.49.0"},"reference-count":46,"publisher":"Springer Science and Business Media LLC","issue":"7","license":[{"start":{"date-parts":[[2025,5,3]],"date-time":"2025-05-03T00:00:00Z","timestamp":1746230400000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2025,5,3]],"date-time":"2025-05-03T00:00:00Z","timestamp":1746230400000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100008766","name":"Fujian Agriculture and Forestry University","doi-asserted-by":"publisher","award":["KFB23157A"],"award-info":[{"award-number":["KFB23157A"]}],"id":[{"id":"10.13039\/501100008766","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/501100001809","name":"the National Natural Science Foundation of China","doi-asserted-by":"crossref","award":["61802064"],"award-info":[{"award-number":["61802064"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"crossref"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["J Supercomput"],"DOI":"10.1007\/s11227-025-07307-6","type":"journal-article","created":{"date-parts":[[2025,5,3]],"date-time":"2025-05-03T10:44:39Z","timestamp":1746269079000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":1,"title":["TMTC: trusted multi-modal transformer classification framework for video frame deletion detection"],"prefix":"10.1007","volume":"81","author":[{"given":"Chunhui","family":"Feng","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Yongxiang","family":"Zhong","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Yigong","family":"Huang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Xiaolong","family":"Liu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2025,5,3]]},"reference":[{"key":"7307_CR1","doi-asserted-by":"crossref","unstructured":"Kumar V, Singh A, Kansal V, et al. (2021) A comprehensive survey on passive video forgery detection techniques\/\/Recent studies on computational intelligence: doctoral symposium on computational intelligence (DoSCI 2020). Springer Singapore, 39-57","DOI":"10.1007\/978-981-15-8469-5_4"},{"issue":"2","key":"7307_CR2","doi-asserted-by":"publisher","first-page":"5415","DOI":"10.1007\/s11042-023-15561-0","volume":"83","author":"NA Shelke","year":"2024","unstructured":"Shelke NA, Kasana SS (2024) Multiple forgery detection in digital video with VGG-16-based deep neural network and KPCA. Multimedia Tools Appl 83(2):5415\u20135435","journal-title":"Multimedia Tools Appl"},{"key":"7307_CR3","doi-asserted-by":"crossref","unstructured":"Xing Q, Luo Y, Zhang Z, et al. (2022) Video inter-frame tampering detection based on SN-VGG+ BiLSTM-AE composite model. In: Proceedings of the 2022 10th international conference on information technology: IoT and smart city. 80\u201387.","DOI":"10.1145\/3582197.3582210"},{"key":"7307_CR4","doi-asserted-by":"crossref","unstructured":"Feng C, Wu D, Wu T, et al. (2024) An MSDCNN-LSTM framework for video frame deletion forensics. Multimedia Tools and Applications, 1\u201320.","DOI":"10.1007\/s11042-024-18324-7"},{"issue":"6","key":"7307_CR5","doi-asserted-by":"publisher","first-page":"865","DOI":"10.1109\/TSMCC.2011.2178594","volume":"42","author":"OP Popoola","year":"2012","unstructured":"Popoola OP, Wang K (2012) Video-based abnormal human behavior recognition\u2014A review. IEEE Trans Syst Man Cybernet Part C (Applications and Reviews) 42(6):865\u2013878","journal-title":"IEEE Trans Syst Man Cybernet Part C (Applications and Reviews)"},{"issue":"3","key":"7307_CR6","first-page":"1","volume":"10","author":"XH Nguyen","year":"2020","unstructured":"Nguyen XH, Hu Y, Amin MA et al (2020) Detecting video inter-frame forgeries based on convolutional neural network model. Int J Image Graph Signal Process 10(3):1","journal-title":"Int J Image Graph Signal Process"},{"key":"7307_CR7","doi-asserted-by":"publisher","first-page":"84","DOI":"10.1016\/j.neucom.2016.03.051","volume":"205","author":"L Yu","year":"2016","unstructured":"Yu L, Wang H, Han Q et al (2016) Exposing frame deletion by detecting abrupt changes in video streams. Neurocomputing 205:84\u201391","journal-title":"Neurocomputing"},{"issue":"2","key":"7307_CR8","doi-asserted-by":"publisher","first-page":"311","DOI":"10.1002\/sec.981","volume":"8","author":"Z Zhang","year":"2015","unstructured":"Zhang Z, Hou J, Ma Q et al (2015) Efficient video frame insertion and deletion detection based on inconsistency of correlations between local binary pattern coded frames. Security Commun Netw 8(2):311\u2013320","journal-title":"Security Commun Netw"},{"issue":"6","key":"7307_CR9","first-page":"91","volume":"21","author":"XJ Yuan","year":"2012","unstructured":"Yuan XJ, Huang TQ, Chen ZW et al (2012) Digital video forgeries detection based on textural features. Comput Syst Appl 21(6):91\u201395","journal-title":"Comput Syst Appl"},{"key":"7307_CR10","doi-asserted-by":"crossref","unstructured":"Bidokhti A, Ghaemmaghami S. (2015) Detection of regional copy\/move forgery in MPEG videos using optical flow. In: 2015 The International Symposium on Artificial Intelligence and Signal Processing (AISP). IEEE, 13\u201317.","DOI":"10.1109\/AISP.2015.7123529"},{"key":"7307_CR11","doi-asserted-by":"crossref","unstructured":"Wang W, Jiang X, Wang S, et al. Identifying video forgery process using optical flow\/\/Digital-Forensics and Watermarking: 12th International Workshop, IWDW 2013, Auckland, New Zealand, October 1\u20134, 2013. Revised Selected Papers 12. Springer Berlin Heidelberg, 2014: 244\u2013257.","DOI":"10.1007\/978-3-662-43886-2_18"},{"key":"7307_CR12","doi-asserted-by":"publisher","first-page":"37196","DOI":"10.1109\/ACCESS.2021.3061586","volume":"9","author":"S Li","year":"2021","unstructured":"Li S, Huo H (2021) Frame deletion detection based on optical flow orientation variation. IEEE Access 9:37196\u201337209","journal-title":"IEEE Access"},{"key":"7307_CR13","doi-asserted-by":"publisher","first-page":"412","DOI":"10.1016\/j.cose.2018.04.013","volume":"77","author":"T Huang","year":"2018","unstructured":"Huang T, Zhang X, Huang W et al (2018) A multi-channel approach through fusion of audio for detecting video inter-frame forgery. Comput Secur 77:412\u2013426","journal-title":"Comput Secur"},{"issue":"12","key":"7307_CR14","doi-asserted-by":"publisher","first-page":"3953","DOI":"10.3390\/s21123953","volume":"21","author":"H Pu","year":"2021","unstructured":"Pu H, Huang T, Weng B et al (2021) Overcome the brightness and jitter noises in video inter-frame tampering detection. Sensors 21(12):3953","journal-title":"Sensors"},{"key":"7307_CR15","doi-asserted-by":"crossref","unstructured":"Long C, Smith E, Basharat A, et al. (2017) A c3d-based convolutional neural network for frame dropping detection in a single video shot. In: 2017 IEEE conference on computer vision and pattern recognition workshops (CVPRW). IEEE, 1898\u20131906.","DOI":"10.1109\/CVPRW.2017.237"},{"key":"7307_CR16","doi-asserted-by":"crossref","unstructured":"Bakas J, Naskar R. (2018) A digital forensic technique for inter\u2013frame video forgery detection based on 3D CNN. In: International conference on information systems security. Cham: Springer International Publishing, 304\u2013317.","DOI":"10.1007\/978-3-030-05171-6_16"},{"issue":"16","key":"7307_CR17","doi-asserted-by":"publisher","first-page":"22731","DOI":"10.1007\/s11042-021-10989-8","volume":"81","author":"NA Shelke","year":"2022","unstructured":"Shelke NA, Kasana SS (2022) Multiple forgery detection and localization technique for digital video using PCT and NBAP. Multimedia Tools Appl 81(16):22731\u201322759","journal-title":"Multimedia Tools Appl"},{"issue":"4","key":"7307_CR18","doi-asserted-by":"publisher","first-page":"3285","DOI":"10.1007\/s11760-023-02990-5","volume":"18","author":"S Li","year":"2024","unstructured":"Li S, Huo H (2024) Continuity-attenuation captured network for frame deletion detection. SIViP 18(4):3285\u20133297","journal-title":"SIViP"},{"key":"7307_CR19","doi-asserted-by":"crossref","unstructured":"Tinipuclla C, Ceron J, Shiguihara P. Frame deletion detection in videos using convolutional neural networks\/\/2024 IEEE ANDESCON. IEEE, 2024: 1\u20136.","DOI":"10.1109\/ANDESCON61840.2024.10755836"},{"key":"7307_CR20","doi-asserted-by":"crossref","unstructured":"Zhang Y, Miao C, Luo M, et al. (2024) MFMS: Learning modality-fused and modality-specific features for deepfake detection and localization tasks. In: Proceedings of the 32nd ACM international conference on multimedia. p 11365\u201311369.","DOI":"10.1145\/3664647.3688984"},{"key":"7307_CR21","doi-asserted-by":"publisher","first-page":"1039","DOI":"10.1109\/TIFS.2022.3233774","volume":"18","author":"C Miao","year":"2023","unstructured":"Miao C, Tan Z, Chu Q et al (2023) F 2 trans: high-frequency fine-grained transformer for face forgery detection. IEEE Trans Inf Forensics Secur 18:1039\u20131051","journal-title":"IEEE Trans Inf Forensics Secur"},{"key":"7307_CR22","first-page":"14200","volume":"34","author":"A Nagrani","year":"2021","unstructured":"Nagrani A, Yang S, Arnab A et al (2021) Attention bottlenecks for multimodal fusion. Adv Neural Inf Process Syst 34:14200\u201314213","journal-title":"Adv Neural Inf Process Syst"},{"key":"7307_CR23","unstructured":"Mohamed A. Deep neural network acoustic models for ASR. University of Toronto, 2014."},{"key":"7307_CR24","doi-asserted-by":"publisher","first-page":"107389","DOI":"10.1016\/j.apacoust.2020.107389","volume":"167","author":"Z Mushtaq","year":"2020","unstructured":"Mushtaq Z, Su SF (2020) Environmental sound classification using a regularized deep convolutional neural network with data augmentation. Appl Acoust 167:107389","journal-title":"Appl Acoust"},{"issue":"2","key":"7307_CR25","doi-asserted-by":"publisher","first-page":"206","DOI":"10.1109\/JSTSP.2019.2908700","volume":"13","author":"H Purwins","year":"2019","unstructured":"Purwins H, Li B, Virtanen T et al (2019) Deep learning for audio signal processing. IEEE J Selected Topics Signal Process 13(2):206\u2013219","journal-title":"IEEE J Selected Topics Signal Process"},{"key":"7307_CR26","doi-asserted-by":"crossref","unstructured":"Xiao X, Zhao S, Zhong X, et al. (2015) A learning-based approach to direction of arrival estimation in noisy and reverberant environments. In: 2015 IEEE international conference on acoustics, speech and signal processing (ICASSP). IEEE, 2814\u20132818.","DOI":"10.1109\/ICASSP.2015.7178484"},{"key":"7307_CR27","doi-asserted-by":"publisher","first-page":"25323","DOI":"10.1109\/ACCESS.2018.2819624","volume":"6","author":"S Jia","year":"2018","unstructured":"Jia S, Xu Z, Wang H et al (2018) Coarse-to-fine copy-move forgery detection for video forensics. IEEE Access 6:25323\u201325335","journal-title":"IEEE Access"},{"key":"7307_CR28","doi-asserted-by":"crossref","unstructured":"Woo S, Park J, Lee J Y, et al. (2018) Cbam: convolutional block attention module. In: Proceedings of the European conference on computer vision (ECCV). 3\u201319.","DOI":"10.1007\/978-3-030-01234-2_1"},{"key":"7307_CR29","unstructured":"Kendall A, Gal Y. What uncertainties do we need in bayesian deep learning for computer vision? Advances in neural information processing systems, 2017, 30."},{"issue":"2","key":"7307_CR30","doi-asserted-by":"publisher","first-page":"2551","DOI":"10.1109\/TPAMI.2022.3171983","volume":"45","author":"Z Han","year":"2022","unstructured":"Han Z, Zhang C, Fu H et al (2022) Trusted multi-view classification with dynamic evidential fusion. IEEE Trans Pattern Anal Mach Intell 45(2):2551\u20132566","journal-title":"IEEE Trans Pattern Anal Mach Intell"},{"key":"7307_CR31","doi-asserted-by":"crossref","unstructured":"Liu W, Luo W, Lian D, et al. (2018) Future frame prediction for anomaly detection\u2013a new baseline. In: Proceedings of the IEEE conference on computer vision and pattern recognition. 6536\u20136545.","DOI":"10.1109\/CVPR.2018.00684"},{"key":"7307_CR32","doi-asserted-by":"crossref","unstructured":"Snoek C G M, Worring M, Smeulders A W M. (2005) Early versus late fusion in semantic video analysis. In: Proceedings of the 13th annual ACM international conference on multimedia. 399\u2013402.","DOI":"10.1145\/1101149.1101236"},{"key":"7307_CR33","doi-asserted-by":"crossref","unstructured":"Tsai Y H H, Bai S, Liang P P, et al. (2019) Multimodal transformer for unaligned multimodal language sequences. In: Proceedings of the conference. Association for Computational Linguistics. Meeting. NIH Public Access, 2019: 6558.","DOI":"10.18653\/v1\/P19-1656"},{"key":"7307_CR34","doi-asserted-by":"crossref","first-page":"1","DOI":"10.1186\/s13635-016-0053-0","volume":"2017","author":"D Shullani","year":"2017","unstructured":"Shullani D, Fontani M, Iuliani M et al (2017) Vision: a video and image dataset for source identification. EURASIP J Inf Secur 2017:1\u201316","journal-title":"EURASIP J Inf Secur"},{"key":"7307_CR35","unstructured":"Soomro K. UCF101: A dataset of 101 human actions classes from videos in the wild. arXiv preprint arXiv:1212.0402, 2012."},{"key":"7307_CR36","doi-asserted-by":"publisher","first-page":"8","DOI":"10.1016\/j.diin.2016.06.003","volume":"18","author":"K Sitara","year":"2016","unstructured":"Sitara K, Mehtre BM (2016) Digital video tampering detection: an overview of passive techniques. Digit Investig 18:8\u201322","journal-title":"Digit Investig"},{"key":"7307_CR37","doi-asserted-by":"crossref","unstructured":"Karpathy A, Toderici G, Shetty S, et al. (2014) Large-scale video classification with convolutional neural networks. In: Proceedings of the IEEE conference on computer vision and pattern recognition. p 1725\u20131732.","DOI":"10.1109\/CVPR.2014.223"},{"key":"7307_CR38","doi-asserted-by":"crossref","unstructured":"Liu Z, Shen Y, Lakshminarasimhan V B, et al. Efficient low-rank multimodal fusion with modality-specific factors. arXiv preprint arXiv:1806.00064, 2018.","DOI":"10.18653\/v1\/P18-1209"},{"key":"7307_CR39","doi-asserted-by":"crossref","unstructured":"Dalal N, Triggs B. (2005) Histograms of oriented gradients for human detection. In: 2005 IEEE computer society conference on computer vision and pattern recognition (CVPR'05). IEEE, 1: 886-893.","DOI":"10.1109\/CVPR.2005.177"},{"key":"7307_CR40","doi-asserted-by":"crossref","unstructured":"Xie S, Sun C, Huang J, et al. (2018) Rethinking spatiotemporal feature learning: Speed-accuracy trade-offs in video classification. In: Proceedings of the European conference on computer vision (ECCV). 305\u2013321.","DOI":"10.1007\/978-3-030-01267-0_19"},{"key":"7307_CR41","doi-asserted-by":"crossref","unstructured":"Carreira J, Zisserman A. (2017) Quo vadis, action recognition? a new model and the kinetics dataset. In: Proceedings of the IEEE conference on computer vision and pattern recognition. 6299\u20136308.","DOI":"10.1109\/CVPR.2017.502"},{"issue":"11","key":"7307_CR42","doi-asserted-by":"publisher","first-page":"301","DOI":"10.3390\/a13110301","volume":"13","author":"G Liu","year":"2020","unstructured":"Liu G, Zhang C, Xu Q et al (2020) I3d-shufflenet based human action recognition. Algorithms 13(11):301","journal-title":"Algorithms"},{"key":"7307_CR43","doi-asserted-by":"crossref","unstructured":"He K, Zhang X, Ren S, et al. (2016) Deep residual learning for image recognition. In: Proceedings of the IEEE conference on computer vision and pattern recognition. p 770\u2013778.","DOI":"10.1109\/CVPR.2016.90"},{"key":"7307_CR44","unstructured":"Simonyan K. Very deep convolutional networks for large-scale image recognition. arXiv preprint arXiv:1409.1556, 2014."},{"key":"7307_CR45","doi-asserted-by":"crossref","unstructured":"Huang G, Liu Z, Van Der Maaten L, et al. (2017) Densely connected convolutional networks. In: Proceedings of the IEEE conference on computer vision and pattern recognition. p 4700\u20134708.","DOI":"10.1109\/CVPR.2017.243"},{"key":"7307_CR46","doi-asserted-by":"crossref","unstructured":"Selvaraju R R, Cogswell M, Das A, et al. (2017) Grad-cam: Visual explanations from deep networks via gradient-based localization. In: Proceedings of the IEEE international conference on computer vision. p 618\u2013626.","DOI":"10.1109\/ICCV.2017.74"}],"container-title":["The Journal of Supercomputing"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11227-025-07307-6.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11227-025-07307-6\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11227-025-07307-6.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,5,3]],"date-time":"2025-05-03T10:44:56Z","timestamp":1746269096000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11227-025-07307-6"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,5,3]]},"references-count":46,"journal-issue":{"issue":"7","published-online":{"date-parts":[[2025,5]]}},"alternative-id":["7307"],"URL":"https:\/\/doi.org\/10.1007\/s11227-025-07307-6","relation":{},"ISSN":["1573-0484"],"issn-type":[{"value":"1573-0484","type":"electronic"}],"subject":[],"published":{"date-parts":[[2025,5,3]]},"assertion":[{"value":"9 April 2025","order":1,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"3 May 2025","order":2,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare no conflict of interest.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}],"article-number":"818"}}