{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,4,10]],"date-time":"2026-04-10T11:23:22Z","timestamp":1775820202444,"version":"3.50.1"},"reference-count":80,"publisher":"Springer Science and Business Media LLC","issue":"1","license":[{"start":{"date-parts":[[2026,4,10]],"date-time":"2026-04-10T00:00:00Z","timestamp":1775779200000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/creativecommons.org\/licenses\/by-nc-nd\/4.0"},{"start":{"date-parts":[[2026,4,10]],"date-time":"2026-04-10T00:00:00Z","timestamp":1775779200000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/creativecommons.org\/licenses\/by-nc-nd\/4.0"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["J. Image Video Process."],"DOI":"10.1186\/s13640-026-00691-w","type":"journal-article","created":{"date-parts":[[2026,4,10]],"date-time":"2026-04-10T10:35:06Z","timestamp":1775817306000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["A review of deep learning-based virtual try-on research"],"prefix":"10.1186","volume":"2026","author":[{"given":"Xiangyan","family":"Fu","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0009-0006-1229-6197","authenticated-orcid":false,"given":"Shengling","family":"Geng","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Weihua","family":"Pu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Xiaojuan","family":"Dang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2026,4,10]]},"reference":[{"key":"691_CR1","unstructured":"M. Bi\u0144kowski, D. Sutherland, M. Arbel, A. Gretton, Demystifying MMD GANs (2018). arXiv:1801.01401"},{"issue":"1","key":"691_CR2","doi-asserted-by":"publisher","first-page":"172","DOI":"10.1109\/TPAMI.2019.2929257","volume":"43","author":"Z Cao","year":"2019","unstructured":"Z. Cao, G. Hidalgo, T. Simon, S. Wei, Y. Sheikh, OpenPose: realtime multi-person 2D pose estimation using part affinity fields. IEEE Trans. Pattern Anal. Mach. Intell. 43(1), 172\u2013186 (2019)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"691_CR3","unstructured":"T. Chang, X. Chen Wei, X. Zhang, Q. Chen, W. Luo, X. Yang, PEMF-VVTO: point-enhanced video virtual try-on via mask-free paradigm (2024). arXiv:2412.03021"},{"key":"691_CR4","doi-asserted-by":"crossref","unstructured":"C. Chen, Y. Chen, H. Shuai, W. Cheng, Size does matter: Size-aware virtual try-on via clothing-oriented transformation try-on network. In Proceedings of the IEEE\/CVF international conference on computer vision, pages 7513\u20137522, (2023)","DOI":"10.1109\/ICCV51070.2023.00691"},{"key":"691_CR5","doi-asserted-by":"crossref","unstructured":"C. Chen, L. Lo, P. Huang, H. Shuai, W. Cheng, Fashionmirror: co-attention feature-remapping virtual try-on with sequential template poses, in Proceedings of the IEEE\/CVF International Conference on Computer Vision (2021), pp. 13809\u201313818","DOI":"10.1109\/ICCV48922.2021.01355"},{"key":"691_CR6","doi-asserted-by":"crossref","unstructured":"S. Choi, S. Park, M. Lee, J. Choo, Viton-HD: high-resolution virtual try-on via misalignment-aware normalization, in Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (2021), pp. 14131\u201314140","DOI":"10.1109\/CVPR46437.2021.01391"},{"key":"691_CR7","unstructured":"Z. Chong, W. Zhang, S. Zhang, J. Zheng, X. Dong, H. Li, Y. Wu, D. Jiang, X. Liang, Catv2ton: Taming Diffusion Transformers for Vision-Based Virtual Try-on with Temporal Concatenation (2025). arXiv:2501.11325"},{"key":"691_CR8","doi-asserted-by":"crossref","unstructured":"A. Chopra, R. Jain, M. Hemani, B. Krishnamurthy, Zflow: gated appearance flow-based virtual try-on with 3D priors, in Proceedings of the IEEE\/CVF International Conference on Computer Vision (2021), pp. 5433\u20135442","DOI":"10.1109\/ICCV48922.2021.00538"},{"key":"691_CR9","doi-asserted-by":"crossref","unstructured":"H. Dong, X. Liang, X. Shen, B. Wang, H. Lai, J. Zhu, Z. Hu, J. Yin, Towards multi-pose guided virtual try-on network, in Proceedings of the IEEE\/CVF International Conference on Computer Vision (2019), pp. 9026\u20139035","DOI":"10.1109\/ICCV.2019.00912"},{"key":"691_CR10","doi-asserted-by":"crossref","unstructured":"H. Dong, X. Liang, X. Shen, B. Wu, B. Chen, J. Yin, FW-GAN: flow-navigated warping GAN for video virtual try-on, in Proceedings of the IEEE\/CVF International Conference on Computer Vision (2019), pp. 1161\u20131170","DOI":"10.1109\/ICCV.2019.00125"},{"key":"691_CR11","doi-asserted-by":"crossref","unstructured":"B. Fele, A. Lampe, P. Peer, V. Struc, C-vton: context-driven image-based virtual try-on network, in Proceedings of the IEEE\/CVF Winter Conference on Applications of Computer Vision (2022), pp. 3144\u20133153","DOI":"10.1109\/WACV51458.2022.00226"},{"key":"691_CR12","doi-asserted-by":"crossref","unstructured":"C. Ge, Y. Song, Y. Ge, H. Yang, W. Liu, P. Luo, Disentangled cycle consistency for highly-realistic virtual try-on, in Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (2021), pp. 16928\u201316937","DOI":"10.1109\/CVPR46437.2021.01665"},{"key":"691_CR13","doi-asserted-by":"crossref","unstructured":"Y. Ge, Y. Song, R. Zhang, C. Ge, W. Liu, P. Luo, Parser-free virtual try-on via distilling appearance flows, in Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (2021), pp. 8485\u20138493","DOI":"10.1109\/CVPR46437.2021.00838"},{"key":"691_CR14","doi-asserted-by":"crossref","unstructured":"J. Gou, S. Sun, J. Zhang, J. Si, C. Qian, L. Zhang, Taming the power of diffusion models for high-quality virtual try-on with appearance flow, in Proceedings of the 31st ACM International Conference on Multimedia (2023), pp. 7599\u20137607","DOI":"10.1145\/3581783.3612255"},{"key":"691_CR15","doi-asserted-by":"crossref","unstructured":"R.A. G\u00fcler, N. Neverova, I. Kokkinos, Densepose: dense human pose estimation in the wild, in Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (2018), pp. 7297\u20137306","DOI":"10.1109\/CVPR.2018.00762"},{"key":"691_CR16","doi-asserted-by":"crossref","unstructured":"X. Han, X. Hu, W. Huang, Clothflow: a flow-based model for clothed person generation, in Proceedings of the IEEE\/CVF International Conference on Computer Vision (2019), pp. 10471\u201310480","DOI":"10.1109\/ICCV.2019.01057"},{"key":"691_CR17","doi-asserted-by":"crossref","unstructured":"X. Han, Z. Wu, Z. Wu, R. Yu, L. Davis, Viton: an image-based virtual try-on network, in Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (2018), pp. 7543\u20137552","DOI":"10.1109\/CVPR.2018.00787"},{"key":"691_CR18","doi-asserted-by":"crossref","unstructured":"X. Han, S. Zheng, Z. Li, C. Wang, X. Sun, Q. Meng, Shape-guided clothing warping for virtual try-on, in Proceedings of the 32nd ACM International Conference on Multimedia (2024), pp. 2593\u20132602","DOI":"10.1145\/3664647.3680756"},{"key":"691_CR19","doi-asserted-by":"crossref","unstructured":"S. He, Y. Song, T. Xiang, Style-based global appearance flow for virtual try-on, in Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (2022), pp. 3470\u20133479","DOI":"10.1109\/CVPR52688.2022.00346"},{"key":"691_CR20","doi-asserted-by":"crossref","unstructured":"Z. He, P. Chen, G. Wang, G. Li, P. Torr, L. Lin, Wildvidfit: video virtual try-on in the wild via image-based controlled diffusion models, in European Conference on Computer Vision (Springer, 2024), pp. 123\u2013139","DOI":"10.1007\/978-3-031-72643-9_8"},{"key":"691_CR21","unstructured":"M. Heusel, H. Ramsauer, T. Unterthiner, B. Nessler, S. Hochreiter, Gans trained by a two time-scale update rule converge to a local Nash equilibrium, in Advances in Neural Information Processing Systems, vol. 30 (2017)"},{"key":"691_CR22","first-page":"6840","volume":"33","author":"J Ho","year":"2020","unstructured":"J. Ho, A. Jain, P. Abbeel, Denoising diffusion probabilistic models. Adv. Neural. Inf. Process. Syst. 33, 6840\u20136851 (2020)","journal-title":"Adv. Neural. Inf. Process. Syst."},{"key":"691_CR23","doi-asserted-by":"crossref","unstructured":"C. Hsieh, C. Chen, C. Chou, Hn. Shuai, J. Liu, W. Cheng, Fashionon: semantic-guided image-based virtual try-on with detailed human and clothing information, in Proceedings of the 27th ACM International Conference on Multimedia (2019), pp. 275\u2013283","DOI":"10.1145\/3343031.3351075"},{"key":"691_CR24","doi-asserted-by":"publisher","first-page":"1233","DOI":"10.1109\/TMM.2022.3143712","volume":"24","author":"B Hu","year":"2022","unstructured":"B. Hu, P. Liu, Z. Zheng, M. Ren, SPG-VTON: semantic prediction guidance for multi-pose virtual try-on. IEEE Trans. Multimed. 24, 1233\u20131246 (2022)","journal-title":"IEEE Trans. Multimed."},{"key":"691_CR25","doi-asserted-by":"crossref","unstructured":"T. Islam, A. Miron, X. Liu, Y. Li, Fashionflow: Leveraging Diffusion Models for Dynamic Fashion Video Synthesis from Static Imagery (2023). arXiv:2310.00106","DOI":"10.3390\/fi16080287"},{"key":"691_CR26","doi-asserted-by":"publisher","DOI":"10.1016\/j.neucom.2024.127887","volume":"594","author":"T Islam","year":"2024","unstructured":"T. Islam, A. Miron, X. Liu, Y. Li, StyleVTON: a multi-pose virtual try-on with identity and clothing detail preservation. Neurocomputing 594, 127887 (2024)","journal-title":"Neurocomputing"},{"key":"691_CR27","doi-asserted-by":"crossref","unstructured":"P. Isola, J. Zhu, T. Zhou, A. Efros, Image-to-image translation with conditional adversarial networks, in Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (2017), pp. 1125\u20131134","DOI":"10.1109\/CVPR.2017.632"},{"key":"691_CR28","doi-asserted-by":"crossref","unstructured":"S. Jandial, A. Chopra, K. Ayush, M. Hemani, B. Krishnamurthy, A. Halwai, Sievenet: a unified framework for robust image-based virtual try-on, in Proceedings of the IEEE\/CVF Winter Conference on Applications of Computer Vision (2020), pp. 2182\u20132190","DOI":"10.1109\/WACV45572.2020.9093458"},{"key":"691_CR29","doi-asserted-by":"crossref","unstructured":"J. Jiang, T. Wang, H. Yan, J. Liu, Clothformer: taming video virtual try-on in all module, in Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (2022), pp. 10799\u201310808","DOI":"10.1109\/CVPR52688.2022.01053"},{"issue":"4","key":"691_CR30","doi-asserted-by":"crossref","first-page":"1","DOI":"10.1145\/3643505","volume":"7","author":"D Kang","year":"2024","unstructured":"D. Kang, E. Baek, S. Son, Y. Lee, T. Gong, H. Kim, Mirror: towards generalizable on-device video virtual try-on for mobile shopping. Proc. ACM Interact. Mob. Wearable Ubiquitous Technol. 7(4), 1\u201327 (2024)","journal-title":"Proc. ACM Interact. Mob. Wearable Ubiquitous Technol."},{"key":"691_CR31","doi-asserted-by":"crossref","unstructured":"J. Kim, G. Gu, M. Park, S. Park, J. Choo, Stableviton: learning semantic correspondence with latent diffusion model for virtual try-on, in Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (2024), pp. 8176\u20138185","DOI":"10.1109\/CVPR52733.2024.00781"},{"key":"691_CR32","doi-asserted-by":"crossref","unstructured":"G. Kuppa, A. Jong, X. Liu, Z. Liu, T. Moh, Shineon: illuminating design choices for practical video-based virtual clothing try-on, in Proceedings of the IEEE\/CVF Winter Conference on Applications of Computer Vision (2021), pp. 191\u2013200","DOI":"10.1109\/WACVW52041.2021.00025"},{"key":"691_CR33","doi-asserted-by":"crossref","unstructured":"H.J. Lee, R. Lee, M. Kang, M. Cho, G. Park, LA-VITON: a network for looking-attractive virtual try-on, in Proceedings of the IEEE\/CVF International Conference on Computer Vision Workshops (2019), pp. 0\u20130","DOI":"10.1109\/ICCVW.2019.00381"},{"key":"691_CR34","doi-asserted-by":"crossref","unstructured":"S. Lee, G. Gu, S. Park, S. Choi, J. Choo, High-resolution virtual try-on with misalignment and occlusion-handled conditions, in European Conference on Computer Vision (Springer, 2022), pp. 204\u2013219","DOI":"10.1007\/978-3-031-19790-1_13"},{"key":"691_CR35","unstructured":"S. Li, Z. Jiang, J. Zhou, Z. Liu, X. Chi, H. Wang, RealVVT: Towards Photorealistic Video Virtual Try-on Via Spatio-temporal Consistency (2025). arXiv:2501.08682"},{"key":"691_CR36","doi-asserted-by":"crossref","unstructured":"Z. Li, P. Wei, X. Yin, Z. Ma, A. Kot, Virtual try-on with pose-garment keypoints guided inpainting, in Proceedings of the IEEE\/CVF International Conference on Computer Vision (2023), pp. 22788\u201322797","DOI":"10.1109\/ICCV51070.2023.02083"},{"issue":"4","key":"691_CR37","doi-asserted-by":"publisher","first-page":"565","DOI":"10.1108\/IJCST-02-2022-0017","volume":"35","author":"W Luo","year":"2023","unstructured":"W. Luo, Y. Zhong, DO-VTON: a details-oriented virtual try-on network. Int. J. Cloth. Sci. Technol. 35(4), 565\u2013580 (2023)","journal-title":"Int. J. Cloth. Sci. Technol."},{"key":"691_CR38","unstructured":"M. Minar, T. Tuan, H. Ahn, P. Rosin, Y. Lai, CP-VTON+: clothing shape and texture preserving image-based virtual try-on, in CVPR Workshops, vol. 3 (2020), pp. 10\u201314"},{"key":"691_CR39","doi-asserted-by":"crossref","unstructured":"Mn. Minar, H. Ahn, Cloth-VTON: clothing three-dimensional reconstruction for hybrid image-based virtual try-on, in Proceedings of the Asian Conference on Computer Vision (2020)","DOI":"10.1007\/978-3-030-69544-6_10"},{"key":"691_CR40","doi-asserted-by":"crossref","unstructured":"D. Morelli, A. Baldrati, G. Cartella, M. Cornia, M. Bertini, R. Cucchiara, LaDI-VTON: latent diffusion textual-inversion enhanced virtual try-on, in Proceedings of the 31st ACM International Conference on Multimedia (2023), pp. 8580\u20138589","DOI":"10.1145\/3581783.3612137"},{"key":"691_CR41","doi-asserted-by":"crossref","unstructured":"D. Morelli, M. Fincato, M. Cornia, F. Landi, F. Cesari, R. Cucchiara, Dress code: high-resolution multi-category virtual try-on, in Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (2022), pp. 2231\u20132235","DOI":"10.1007\/978-3-031-20074-8_20"},{"key":"691_CR42","doi-asserted-by":"crossref","unstructured":"H. Nguyen, Q. Nguyen, K. Nguyen, R. Nguyen, SwiftTry: fast and consistent video virtual try-on with diffusion models, in Proceedings of the AAAI Conference on Artificial Intelligence (2025), pp. 6200\u20136208","DOI":"10.1609\/aaai.v39i6.32663"},{"key":"691_CR43","doi-asserted-by":"crossref","unstructured":"X. Nie, J. Feng, J. Zhang, S. Yan, Single-stage multi-person pose machines, in Proceedings of the IEEE\/CVF International Conference on Computer Vision (2019), pp. 6951\u20136960","DOI":"10.1109\/ICCV.2019.00705"},{"key":"691_CR44","unstructured":"A. Radford, J. Wook, H. Aditya, R. Gabriel, G. Sandhini, G. Sastry, A. Askell, P. Mishkin, J. Clark, G. Krueger et al., Clip: Learning Transferable Visual Models from Natural Language Supervision (2019)"},{"key":"691_CR45","doi-asserted-by":"crossref","unstructured":"R. Rombach, A. Blattmann, D. Lorenz, P. Esser, B. Ommer, High-resolution image synthesis with latent diffusion models, in Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (2022), pp. 10684\u201310695","DOI":"10.1109\/CVPR52688.2022.01042"},{"key":"691_CR46","doi-asserted-by":"crossref","unstructured":"O. Ronneberger, P. Fischer, T. Brox, U-net: convolutional networks for biomedical image segmentation, in Medical Image Computing and Computer-Assisted Intervention-MICCAI 2015: 18th International Conference, Munich, Germany, October 5\u20139, 2015, Proceedings, part III 18 (Springer, 2015), pp. 234\u2013241","DOI":"10.1007\/978-3-319-24574-4_28"},{"key":"691_CR47","unstructured":"I. Santesteban Garay, Data-Driven Models of 3D Avatars and Clothing for Virtual Try-On. Ph.D. Thesis, Universidad Rey Juan Carlos (2022)"},{"key":"691_CR48","doi-asserted-by":"crossref","unstructured":"S. Shim, J. Chung, J. Heo, Towards squeezing-averse virtual try-on via sequential deformation, in Proceedings of the AAAI Conference on Artificial Intelligence (2024), pp. 4856\u20134863","DOI":"10.1609\/aaai.v38i5.28288"},{"key":"691_CR49","unstructured":"J. Song, C. Meng, S. Ermon, Denoising Diffusion Implicit Models (2020). arXiv:2010.02502"},{"key":"691_CR50","doi-asserted-by":"crossref","unstructured":"K. Sun, B. Xiao, D. Liu, J. Wang, Deep high-resolution representation learning for human pose estimation, in Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (2019), pp. 5693\u20135703","DOI":"10.1109\/CVPR.2019.00584"},{"key":"691_CR51","doi-asserted-by":"crossref","unstructured":"M. Tran, J. Clements, A. Prasanna Manoharan, T. Nguyen, N. Le, Dualfit: a two-stage virtual try-on via warping and synthesis, in Proceedings of the IEEE\/CVF International Conference on Computer Vision (2025), pp. 2397\u20132407","DOI":"10.1109\/ICCVW69036.2025.00252"},{"key":"691_CR52","doi-asserted-by":"publisher","first-page":"114367","DOI":"10.1109\/ACCESS.2021.3104274","volume":"9","author":"T Tuan","year":"2021","unstructured":"T. Tuan, M. Minar, H. Ahn, J. Wainwright, Multiple pose virtual try-on based on 3D clothing reconstruction. IEEE Access 9, 114367\u2013114380 (2021)","journal-title":"IEEE Access"},{"issue":"12","key":"691_CR53","doi-asserted-by":"publisher","first-page":"5255","DOI":"10.3390\/app14125255","volume":"14","author":"Y Wan","year":"2024","unstructured":"Y. Wan, N. Ding, L. Yao, FA-VTON: a feature alignment-based model for virtual try-on. Appl. Sci. 14(12), 5255 (2024)","journal-title":"Appl. Sci."},{"key":"691_CR54","doi-asserted-by":"crossref","unstructured":"B. Wang, H. Zheng, X. Liang, Y. Chen, L. Lin, M. Yang, Toward characteristic-preserving image-based virtual try-on network, in Proceedings of the European Conference on Computer Vision (ECCV) (2018), pp. 589\u2013604","DOI":"10.1007\/978-3-030-01261-8_36"},{"key":"691_CR55","doi-asserted-by":"crossref","unstructured":"H. Wang, Z. Zhang, D. Di, S. Zhang, W. Zuo, MV-VTON: multi-view virtual try-on with diffusion models, in Proceedings of the AAAI Conference on Artificial Intelligence (2025), pp. 7682\u20137690","DOI":"10.1609\/aaai.v39i7.32827"},{"key":"691_CR56","doi-asserted-by":"crossref","unstructured":"Y. Wang, W. Dai, L. Chan, H. Zhou, A. Zhang, S. Liu, GPD-VVTO: preserving garment details in video virtual try-on, in Proceedings of the 32nd ACM International Conference on Multimedia (2024), pp. 7133\u20137142","DOI":"10.1145\/3664647.3680701"},{"issue":"4","key":"691_CR57","doi-asserted-by":"publisher","first-page":"600","DOI":"10.1109\/TIP.2003.819861","volume":"13","author":"Z Wang","year":"2004","unstructured":"Z. Wang, A. Bovik, H. Sheikh, E. Simoncelli, Image quality assessment: from error visibility to structural similarity. IEEE Trans. Image Process. 13(4), 600\u2013612 (2004)","journal-title":"IEEE Trans. Image Process."},{"key":"691_CR58","doi-asserted-by":"crossref","unstructured":"Z. Xie, Z. Huang, X. Dong, F. Zhao, H. Dong, X. Zhang, F. Zhu, X, Liang, GP-VTON: towards general purpose virtual try-on via collaborative local-flow global-parsing learning, in Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (2023), pp. 23550\u201323559","DOI":"10.1109\/CVPR52729.2023.02255"},{"key":"691_CR59","doi-asserted-by":"crossref","unstructured":"Z. Xie, J. Lai, X. Xie, LG-VTON: fashion landmark meets image-based virtual try-on, in Chinese Conference on Pattern Recognition and Computer Vision (PRCV) (Springer, 2020), pp. 286\u2013297","DOI":"10.1007\/978-3-030-60636-7_24"},{"key":"691_CR60","doi-asserted-by":"crossref","unstructured":"Y. Xu, T. Gu, W. Chen, A. Chen, OOTDiffusion: outfitting fusion based latent diffusion for controllable virtual try-on, in Proceedings of the AAAI Conference on Artificial Intelligence (2025), pp. 8996\u20139004","DOI":"10.1609\/aaai.v39i9.32973"},{"key":"691_CR61","doi-asserted-by":"crossref","unstructured":"Z. Xu, M. Chen, Z. Wang, L. Xing, Z. Zhai, N. Sang, J. Lan, S. Xiao, C. Gao, Tunnel try-on: excavating spatial-temporal tunnels for high-quality virtual try-on in videos, in Proceedings of the 32nd ACM International Conference on Multimedia (2024), pp. 3199\u20133208","DOI":"10.1145\/3664647.3680836"},{"key":"691_CR62","doi-asserted-by":"crossref","unstructured":"K. Yan, T. Gao, H. Zhang, C. Xie, Linking garment with person via semantically associated landmarks for virtual try-on, in Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (2023), pp. 17194\u201317204","DOI":"10.1109\/CVPR52729.2023.01649"},{"key":"691_CR63","doi-asserted-by":"crossref","unstructured":"H. Yang, X. Yu, Z. Liu, Full-range virtual try-on with recurrent tri-level transform, in Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (2022), pp. 3460\u20133469","DOI":"10.1109\/CVPR52688.2022.00345"},{"key":"691_CR64","doi-asserted-by":"crossref","unstructured":"H. Yang, R. Zhang, X. Guo, W. Liu, W. Zuo, P. Luo, Towards photo-realistic virtual try-on by adaptively generating-preserving image content, in Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (2020), pp. 7850\u20137859","DOI":"10.1109\/CVPR42600.2020.00787"},{"key":"691_CR65","doi-asserted-by":"crossref","unstructured":"X. Yang, C. Ding, Z. Hong, J. Huang, J. Tao, X. Xu, Texture-preserving diffusion models for high-fidelity virtual try-on, in Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (2024), pp. 7017\u20137026","DOI":"10.1109\/CVPR52733.2024.00670"},{"key":"691_CR66","doi-asserted-by":"publisher","first-page":"1477","DOI":"10.1109\/TMM.2023.3234399","volume":"25","author":"Z Yang","year":"2023","unstructured":"Z. Yang, J. Chen, Y. Shi, H. Li, T. Chen, L. Lin, OccluMix: towards de-occlusion virtual try-on by semantically-guided mixup. IEEE Trans. Multimed. 25, 1477\u20131488 (2023)","journal-title":"IEEE Trans. Multimed."},{"issue":"5","key":"691_CR67","doi-asserted-by":"publisher","first-page":"3297","DOI":"10.1007\/s00371-024-03603-z","volume":"41","author":"J Ye","year":"2025","unstructured":"J. Ye, Y. Wang, F. Xie, Q. Wang, X. Gu, Z. Wu, Slot-VTON: subject-driven diffusion-based virtual try-on with slot attention. Vis. Comput. 41(5), 3297\u20133308 (2025)","journal-title":"Vis. Comput."},{"issue":"4","key":"691_CR68","doi-asserted-by":"publisher","first-page":"1101","DOI":"10.1109\/TCE.2023.3306206","volume":"69","author":"F Yu","year":"2023","unstructured":"F. Yu, A. Hua, C. Du, M. Jiang, X. Wei, T. Peng, L. Xu, X. Hu, VTON-MP: multi-pose virtual try-on via appearance flow and feature filtering. IEEE Trans. Consum. Electron. 69(4), 1101\u20131113 (2023)","journal-title":"IEEE Trans. Consum. Electron."},{"key":"691_CR69","doi-asserted-by":"crossref","unstructured":"R. Yu, X. Wang, X. Xie, VTNFP: an image-based virtual try-on network with body and clothing feature preservation, in Proceedings of the IEEE\/CVF International Conference on Computer Vision (2019), pp. 10511\u201310520","DOI":"10.1109\/ICCV.2019.01061"},{"key":"691_CR70","doi-asserted-by":"crossref","unstructured":"L. Zhang, A. Rao, M. Agrawala, Adding conditional control to text-to-image diffusion models, in Proceedings of the IEEE\/CVF International Conference on Computer Vision (2023), pp. 3836\u20133847","DOI":"10.1109\/ICCV51070.2023.00355"},{"key":"691_CR71","doi-asserted-by":"crossref","unstructured":"R. Zhang, P. Isola, A. Efros, E. Shechtman, O. Wang, The unreasonable effectiveness of deep features as a perceptual metric, in Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (2018), pp. 586\u2013595","DOI":"10.1109\/CVPR.2018.00068"},{"key":"691_CR72","unstructured":"X. Zhang, E. Lin, X. Li, Y. Luo, M. Kampffmeyer, X. Dong, X. Liang, MMTryon: Multi-modal Multi-reference Control for High-Quality Fashion Generation (2024). arXiv:2405.00448"},{"key":"691_CR73","doi-asserted-by":"crossref","unstructured":"F. Zhao, Z. Xie, M. Kampffmeyer, H. Dong, S. Han, T. Zheng, T. Zhang, X. Liang, M3D-VTON: a monocular-to-3D virtual try-on network, in Proceedings of the IEEE\/CVF International Conference on Computer Vision (2021), pp. 13239\u201313249","DOI":"10.1109\/ICCV48922.2021.01299"},{"key":"691_CR74","unstructured":"J. Zheng, J. Wang, F. Zhao, X. Zhang, X. Liang, Dynamic Try-on: Taming Video Virtual Try-on with Dynamic Attention Mechanism (2024). arXiv:2412.09822"},{"key":"691_CR75","unstructured":"J. Zheng, F. Zhao, Y. Xu, X. Dong, X. Liang, VITON-DiT: Learning In-the-Wild Video Try-on from Human Dance Videos Via Diffusion Transformers (2024). arXiv:2405.18326"},{"key":"691_CR76","doi-asserted-by":"crossref","unstructured":"X. Zhong, Z. Wu, T. Tan, G. Lin, Q. Wu, MV-TON: memory-based video virtual try-on network, in Proceedings of the 29th ACM International Conference on Multimedia (2021), pp. 908\u2013916","DOI":"10.1145\/3474085.3475269"},{"key":"691_CR77","doi-asserted-by":"crossref","unstructured":"X. Zhong, Z. Wu, X. Yang, G. Lin, Q. Wu, IPVTON: Image-Based 3D Virtual Try-on with Image Prompt Adapter (2025). arXiv:2501.15616","DOI":"10.1609\/aaai.v39i10.33159"},{"key":"691_CR78","doi-asserted-by":"crossref","unstructured":"J. Zhu, T. Park, P. Isola, A. Efros, Unpaired image-to-image translation using cycle-consistent adversarial networks, in Proceedings of the IEEE International Conference on Computer Vision (2017), pp. 2223\u20132232","DOI":"10.1109\/ICCV.2017.244"},{"key":"691_CR79","doi-asserted-by":"crossref","unstructured":"L. Zhu, D. Yang, T. Zhu, F. Reda, W. Chan, C. Saharia, M. Norouzi, I. Kemelmacher-Shlizerman, TryOnDiffusion: a tale of two unets, in Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (2023), pp. 4606\u20134615","DOI":"10.1109\/CVPR52729.2023.00447"},{"key":"691_CR80","unstructured":"T. Zuo, Z. Huang, S. Ning, E. Lin, C. Liang, Z. Zheng, DreamVVT: Mastering Realistic Video Virtual Try-on in the Wild Via a Stage-Wise Diffusion Transformer Framework (2025). arXiv:2508.02807"}],"container-title":["Journal on Image and Video Processing"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1186\/s13640-026-00691-w.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1186\/s13640-026-00691-w","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1186\/s13640-026-00691-w.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,4,10]],"date-time":"2026-04-10T10:35:52Z","timestamp":1775817352000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1186\/s13640-026-00691-w"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,4,10]]},"references-count":80,"journal-issue":{"issue":"1","published-online":{"date-parts":[[2026,12]]}},"alternative-id":["691"],"URL":"https:\/\/doi.org\/10.1186\/s13640-026-00691-w","relation":{},"ISSN":["3091-454X"],"issn-type":[{"value":"3091-454X","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,4,10]]},"assertion":[{"value":"14 July 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"22 February 2026","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"10 April 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare no competing interests.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Competing interests"}}],"article-number":"5"}}