{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,3,22]],"date-time":"2026-03-22T03:03:20Z","timestamp":1774148600883,"version":"3.50.1"},"reference-count":30,"publisher":"Springer Science and Business Media LLC","issue":"8","license":[{"start":{"date-parts":[[2024,3,19]],"date-time":"2024-03-19T00:00:00Z","timestamp":1710806400000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2024,3,19]],"date-time":"2024-03-19T00:00:00Z","timestamp":1710806400000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"name":"the STI 2030-Major Projects of China","award":["2021ZD0201300"],"award-info":[{"award-number":["2021ZD0201300"]}]},{"name":"the National Science Foundation of China","award":["62276127"],"award-info":[{"award-number":["62276127"]}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Multimed Tools Appl"],"DOI":"10.1007\/s11042-024-18868-8","type":"journal-article","created":{"date-parts":[[2024,3,19]],"date-time":"2024-03-19T09:02:22Z","timestamp":1710838942000},"page":"4343-4359","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":6,"title":["RandoMix: a mixed sample data augmentation method with multiple mixed modes"],"prefix":"10.1007","volume":"84","author":[{"given":"Xiaoliang","family":"Liu","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-7285-326X","authenticated-orcid":false,"given":"Furao","family":"Shen","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Jian","family":"Zhao","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Changhai","family":"Nie","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2024,3,19]]},"reference":[{"key":"18868_CR1","doi-asserted-by":"crossref","unstructured":"Voulodimos A, Doulamis N, Doulamis A, Protopapadakis E (2018) Deep learning for computer vision: a brief review. Computational intelligence and neuroscience","DOI":"10.1155\/2018\/7068349"},{"key":"18868_CR2","unstructured":"Kamath U, Liu J, Whitaker J Deep learning for NLP and speech recognition vol 84. Springer"},{"key":"18868_CR3","first-page":"781","volume":"181","author":"V Vapnik","year":"1968","unstructured":"Vapnik V (1968) On the uniform convergence of relative frequencies of events to their probabilities. Dokl Akad Nauk USSR 181:781\u2013787","journal-title":"Dokl Akad Nauk USSR"},{"key":"18868_CR4","unstructured":"Zhang H, Cisse M, Dauphin YN, Lopez-Paz D (2018) Mixup: beyond empirical risk minimization. International conference on learning representations (ICLR)"},{"key":"18868_CR5","doi-asserted-by":"crossref","unstructured":"Yun S, Han D, Oh SJ, Chun S, Choe J, Yoo Y (2019) Cutmix: regularization strategy to train strong classifiers with localizable features. IEEE international conference on computer vision (ICCV)","DOI":"10.1109\/ICCV.2019.00612"},{"key":"18868_CR6","unstructured":"Qin J, Fang J, Zhang Q, Liu W, Wang X, Wang X (2020) Resizemix: mixing data with preserved object information and true labels. arXiv:2012.11101"},{"key":"18868_CR7","unstructured":"Harris E, Marcu A, Painter M, Niranjan M, Hare AP-BJ (2021) Fmix: enhancing mixed sample data augmentation. International conference on learning representations (ICLR)"},{"key":"18868_CR8","unstructured":"Kim J-H, Choo W, Song HO (2020) Puzzle mix: exploiting saliency and local statistics for optimal mixup. International conference on machine learning (ICML)"},{"key":"18868_CR9","unstructured":"Kim J, Choo W, Jeong H, Song HO (2021) Co-mixup: saliency guided joint mixup with supermodular diversity. In: international conference on learning representations (ICLR)"},{"key":"18868_CR10","unstructured":"Uddin AFMS, Monira MS, Shin W, Chung T, Bae S-H (2021) Saliencymix: a saliency guided data augmentation strategy for better regularization. International conference on learning representations (ICLR)"},{"key":"18868_CR11","unstructured":"Krizhevsky A, Hinton G et al (2009) Learning multiple layers of features from tiny images"},{"key":"18868_CR12","unstructured":"Chrabaszcz P, Loshchilov I, Hutter F (2017) A downsampled variant of imagenet as an alternative to the cifar datasets. arXiv:1707.08819"},{"key":"18868_CR13","doi-asserted-by":"crossref","unstructured":"Russakovsky O, Deng J, Su H, Krause J, Satheesh S, Ma S, Huang Z, Karpathy A, Khosla A, Bernstein M, et al (2015) Imagenet large scale visual recognition challenge. International journal on computer vision (IJCV)","DOI":"10.1007\/s11263-015-0816-y"},{"key":"18868_CR14","unstructured":"Warden P (2017) Speech commands: a public dataset for single-word speech recognition. Dataset available from http:\/\/download.tensorflow.org\/data\/speech_commands_v0.01.tar.gz"},{"key":"18868_CR15","unstructured":"Verma V, Lamb A, Beckham C, Najafi A, Mitliagkas I, Courville A, Lopez-Paz D, Bengio Y (2019) Manifold mixup: better representations by interpolating hidden states. International conference on machine learning (ICML)"},{"key":"18868_CR16","first-page":"589","volume":"36","author":"M Faramarzi","year":"2022","unstructured":"Faramarzi M, Amini M, Badrinaaraayanan A, Verma V, Chandar S (2022) Patchup: a feature-space block-level regularization technique for convolutional neural networks. Proc AAAI Conf Artif Intell 36:589\u2013597","journal-title":"Proc AAAI Conf Artif Intell"},{"key":"18868_CR17","unstructured":"DeVries T, Taylor GW (2017) Improved regularization of convolutional neural networks with cutout. arXiv:1708.04552"},{"key":"18868_CR18","doi-asserted-by":"crossref","unstructured":"He K, Zhang X, Ren S, Sun J (2016) Identity mappings in deep residual networks. In: European conference on computer vision (ECCV)","DOI":"10.1007\/978-3-319-46493-0_38"},{"key":"18868_CR19","doi-asserted-by":"crossref","unstructured":"Zagoruyko S, Komodakis N (2016) Wide residual networks. In: Procedings of the British machine vision conference 2016. British machine vision association","DOI":"10.5244\/C.30.87"},{"key":"18868_CR20","doi-asserted-by":"crossref","unstructured":"He K, Zhang X, Ren S, Sun J (2016) Deep residual learning for image recognition. In: IEEE conference on computer vision and pattern recognition (CVPR)","DOI":"10.1109\/CVPR.2016.90"},{"key":"18868_CR21","doi-asserted-by":"crossref","unstructured":"Liu Z, Lin Y, Cao Y, Hu H, Wei Y, Zhang Z, Lin S, Guo B (2021) Swin transformer: hierarchical vision transformer using shifted windows. In: IEEE conference on computer vision and pattern recognition (CVPR), pp 10012\u201310022","DOI":"10.1109\/ICCV48922.2021.00986"},{"key":"18868_CR22","unstructured":"Loshchilov I, Hutter F (2017) Sgdr: stochastic gradient descent with warm restarts. In: International conference on learning representations (ICLR)"},{"key":"18868_CR23","first-page":"12826","volume":"35","author":"W Wang","year":"2022","unstructured":"Wang W, Liang J, Liu D (2022) Learning equivariant segmentation with instance-unique querying. Neural Inf Process Syst (NeurIPS) 35:12826\u201312840","journal-title":"Neural Inf Process Syst (NeurIPS)"},{"key":"18868_CR24","unstructured":"Wang W, Han C, Zhou T, Liu D (2023) Visual recognition with deep nearest centroids. In: International conference on learning representations (ICLR)"},{"key":"18868_CR25","doi-asserted-by":"crossref","unstructured":"Liu D, Cui Y, Tan W, Chen Y (2021) Sg-net: spatial granularity network for one-stage video instance segmentation. In: IEEE conference on computer vision and pattern recognition (CVPR), pp 9816\u20139825","DOI":"10.1109\/CVPR46437.2021.00969"},{"key":"18868_CR26","unstructured":"Liang J, Zhou T, Liu D, Wang W (2023) Clustseg: clustering for universal segmentation. In: international conference on machine learning (ICML)"},{"key":"18868_CR27","doi-asserted-by":"crossref","unstructured":"Selvaraju RR, Cogswell M, Das A, Vedantam R, Parikh D, Batra D (2017) Grad-cam: visual explanations from deep networks via gradient-based localization. In: IEEE international conference on computer vision (ICCV), pp 618\u2013626","DOI":"10.1109\/ICCV.2017.74"},{"key":"18868_CR28","unstructured":"Brain G (2017) Tensorflow speech recognition challenge. https:\/\/www.kaggle.com\/c\/tensorflow-speech-recognition-challenge"},{"key":"18868_CR29","unstructured":"Goodfellow IJ, Shlens J, Szegedy C (2015) Explaining and harnessing adversarial examples. International conference on learning representations (ICLR)"},{"key":"18868_CR30","unstructured":"Hendrycks D, Dietterich T (2019) Benchmarking neural network robustness to common corruptions and perturbations. International conference on learning representations (ICLR)"}],"container-title":["Multimedia Tools and Applications"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11042-024-18868-8.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11042-024-18868-8\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11042-024-18868-8.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,3,23]],"date-time":"2025-03-23T00:14:28Z","timestamp":1742688868000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11042-024-18868-8"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2024,3,19]]},"references-count":30,"journal-issue":{"issue":"8","published-online":{"date-parts":[[2025,3]]}},"alternative-id":["18868"],"URL":"https:\/\/doi.org\/10.1007\/s11042-024-18868-8","relation":{},"ISSN":["1573-7721"],"issn-type":[{"value":"1573-7721","type":"electronic"}],"subject":[],"published":{"date-parts":[[2024,3,19]]},"assertion":[{"value":"23 October 2023","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"2 February 2024","order":2,"name":"revised","label":"Revised","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"4 March 2024","order":3,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"19 March 2024","order":4,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no conflict of interest.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflicts of interest"}}]}}