{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,6,2]],"date-time":"2026-06-02T16:01:50Z","timestamp":1780416110201,"version":"3.54.1"},"reference-count":90,"publisher":"Springer Science and Business Media LLC","issue":"3","license":[{"start":{"date-parts":[[2026,5,11]],"date-time":"2026-05-11T00:00:00Z","timestamp":1778457600000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2026,5,11]],"date-time":"2026-05-11T00:00:00Z","timestamp":1778457600000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Mach. Intell. Res."],"published-print":{"date-parts":[[2026,6]]},"DOI":"10.1007\/s11633-025-1623-x","type":"journal-article","created":{"date-parts":[[2026,5,11]],"date-time":"2026-05-11T03:51:04Z","timestamp":1778471464000},"page":"545-564","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["Deep Learning-based Face Video Restoration Technique: A Survey"],"prefix":"10.1007","volume":"23","author":[{"ORCID":"https:\/\/orcid.org\/0000-0002-1156-9193","authenticated-orcid":false,"given":"Chenyang","family":"Wang","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Kangmeng","family":"Yu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Kui","family":"Jiang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Xianming","family":"Liu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-5694-505X","authenticated-orcid":false,"given":"Junjun","family":"Jiang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2026,5,11]]},"reference":[{"issue":"8","key":"1623_CR1","doi-asserted-by":"publisher","first-page":"2244","DOI":"10.1109\/TCSVT.2018.2868063","volume":"29","author":"Y Zhang","year":"2019","unstructured":"Y. Zhang, X. Gao, L. He, W. Lu, R. He. Blind video quality assessment with weakly supervised learning and resampling strategy. IEEE Transactions on Circuits and Systems for Video Technology, vol. 29, no. 8, pp. 2244\u20132255, 2019. DOI: https:\/\/doi.org\/10.1109\/TCSVT.2018.2868063.","journal-title":"IEEE Transactions on Circuits and Systems for Video Technology"},{"issue":"9","key":"1623_CR2","doi-asserted-by":"publisher","first-page":"2008","DOI":"10.1109\/JPROC.2013.2257632","volume":"101","author":"A C Bovik","year":"2013","unstructured":"A. C. Bovik. Automatic prediction of perceptual image and video quality. Proceedings of the IEEE, vol. 101, no. 9, pp. 2008\u20132024, 2013. DOI: https:\/\/doi.org\/10.1109\/JPROC.2013.2257632.","journal-title":"Proceedings of the IEEE"},{"key":"1623_CR3","doi-asserted-by":"publisher","first-page":"5929","DOI":"10.1109\/CVPRW63382.2024.00600","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition Workshops","author":"Z Chen","year":"2024","unstructured":"Z. Chen, J. He, X. Lin, Y. Qiao, C. Dong. Towards real-world video face restoration: A new benchmark. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition Workshops, Seattle, USA, pp. 5929\u20135939, 2024. DOI: https:\/\/doi.org\/10.1109\/CVPRW63382.2024.00600."},{"issue":"12","key":"1623_CR4","doi-asserted-by":"publisher","first-page":"11892","DOI":"10.1109\/TPAMI.2025.3598132","volume":"47","author":"J Jiang","year":"2025","unstructured":"J. Jiang, Z. Zuo, G. Wu, K. Jiang, X. Liu. A survey on all-in-one image restoration: Taxonomy, evaluation and future trends. IEEE Transactions on Pattern Analysis and Machine Intelligence, vol. 47, no. 12, pp. 11892\u201311911, 2025. DOI: https:\/\/doi.org\/10.1109\/TPAMI.2025.3598132.","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"key":"1623_CR5","doi-asserted-by":"publisher","first-page":"9164","DOI":"10.1109\/CVPR46437.2021.00905","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"X Wang","year":"2021","unstructured":"X. Wang, Y. Li, H. Zhang, Y. Shan. Towards real-world blind face restoration with generative facial prior. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, Nashville, USA, pp. 9164\u20139174, 2021. DOI: https:\/\/doi.org\/10.1109\/CVPR46437.2021.00905."},{"key":"1623_CR6","doi-asserted-by":"publisher","first-page":"1551","DOI":"10.1145\/3394171.3413965","volume-title":"Proceedings of the 28th ACM International Conference on Multimedia","author":"L Yang","year":"2020","unstructured":"L. Yang, S. Wang, S. Ma, W. Gao, C. Liu, P. Wang, P. Ren. HiFaceGAN: Face renovation via collaborative suppression and replenishment. In Proceedings of the 28th ACM International Conference on Multimedia, Seattle, USA, pp. 1551\u20131560, 2020. DOI: https:\/\/doi.org\/10.1145\/3394171.3413965."},{"key":"1623_CR7","volume-title":"Proceedings of the 37th International Conference on Neural Information Processing Systems","author":"P Yang","year":"2023","unstructured":"P. Yang, S. Zhou, Q. Tao, C. C. Loy. PGDiff: Guiding diffusion models for versatile face restoration via partial guidance. In Proceedings of the 37th International Conference on Neural Information Processing Systems, New Orleans, USA, Article number 1398, 2023."},{"issue":"12","key":"1623_CR8","doi-asserted-by":"publisher","first-page":"9991","DOI":"10.1109\/TPAMI.2024.3432651","volume":"46","author":"Z Yue","year":"2024","unstructured":"Z. Yue, C. C. Loy. DifFace: Blind face restoration with diffused error contraction. IEEE Transactions on Pattern Analysis and Machine Intelligence, vol. 46, no. 12, pp. 9991\u201310004, 2024. DOI: https:\/\/doi.org\/10.1109\/TPAMI.2024.3432651.","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"issue":"11","key":"1623_CR9","doi-asserted-by":"publisher","first-page":"139","DOI":"10.1145\/3422622","volume":"63","author":"I Goodfellow","year":"2020","unstructured":"I. Goodfellow, J. Pouget-Abadie, M. Mirza, B. Xu, D. Warde-Farley, S. Ozair, A. Courville, Y. Bengio. Generative adversarial networks. Communications of the ACM, vol. 63, no. 11, pp. 139\u2013144, 2020. DOI: https:\/\/doi.org\/10.1145\/3422622.","journal-title":"Communications of the ACM"},{"key":"1623_CR10","first-page":"2256","volume-title":"Proceedings of the 32nd International Conference on International Conference on Machine Learning","author":"J Sohl-Dickstein","year":"2015","unstructured":"J. Sohl-Dickstein, E. A. Weiss, N. Maheswaranathan, S. Ganguli. Deep unsupervised learning using nonequilibrium thermodynamics. In Proceedings of the 32nd International Conference on International Conference on Machine Learning, Lille, France, pp. 2256\u20132265, 2015."},{"key":"1623_CR11","volume-title":"Proceedings of the 36th International Conference on Neural Information Processing Systems","author":"S Zhou","year":"2022","unstructured":"S. Zhou, K. C. K. Chan, C. Li, C. C. Loy. Towards robust blind face restoration with codebook lookup transformer. In Proceedings of the 36th International Conference on Neural Information Processing Systems, New Orleans, USA, Article number 2218, 2022."},{"key":"1623_CR12","doi-asserted-by":"publisher","first-page":"126","DOI":"10.1007\/978-3-031-19797-0_8","volume-title":"Proceedings of the 17th European Conference on Computer Vision","author":"Y Gu","year":"2022","unstructured":"Y. Gu, X. Wang, L. Xie, C. Dong, G. Li, Y. Shan, M. M. Cheng. VQFR: Blind face restoration with vector-quantized dictionary and parallel decoder. In Proceedings of the 17th European Conference on Computer Vision, Tel Aviv, Israel, pp. 126\u2013143, 2022. DOI: https:\/\/doi.org\/10.1007\/978-3-031-19797-0_8."},{"key":"1623_CR13","doi-asserted-by":"publisher","first-page":"17491","DOI":"10.1109\/CVPR52688.2022.01699","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"Z Wang","year":"2022","unstructured":"Z. Wang, J. Zhang, R. Chen, W. Wang, P. Luo. Restore-Former: High-quality blind face restoration from undegraded key-value pairs. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, New Orleans, USA, pp. 17491\u201317500, 2022. DOI: https:\/\/doi.org\/10.1109\/CVPR52688.2022.01699."},{"key":"1623_CR14","doi-asserted-by":"publisher","first-page":"1306","DOI":"10.24963\/ijcai.2022\/182","volume-title":"Proceedings of the 31st International Joint Conference on Artificial Intelligence","author":"J Shi","year":"2022","unstructured":"J. Shi, Y. Wang, S. Dong, X. Hong, Z. Yu, F. Wang, C. Wang, Y. Gong. IDPT: Interconnected dual pyramid transformer for face super-resolution. In Proceedings of the 31st International Joint Conference on Artificial Intelligence, Vienna, Austria, pp. 1306\u20131312, 2022. DOI: https:\/\/doi.org\/10.24963\/ijcai.2022\/182."},{"key":"1623_CR15","doi-asserted-by":"publisher","first-page":"1978","DOI":"10.1109\/TIP.2023.3261747","volume":"32","author":"G Gao","year":"2023","unstructured":"G. Gao, Z. Xu, J. Li, J. Yang, T. Zeng, G. J. Qi. CTCNet: A CNN-transformer cooperation network for face image super-resolution. IEEE Transactions on Image Processing, vol. 32, pp. 1978\u20131991, 2023. DOI: https:\/\/doi.org\/10.1109\/TIP.2023.3261747.","journal-title":"IEEE Transactions on Image Processing"},{"key":"1623_CR16","doi-asserted-by":"publisher","first-page":"4515","DOI":"10.1145\/3664647.3681088","volume-title":"Proceedings of the 32nd ACM International Conference on Multimedia","author":"W Li","year":"2024","unstructured":"W. Li, H. Guo, X. Liu, K. Liang, J. Hu, Z. Ma, J. Guo. Efficient face super-resolution via wavelet-based feature enhancement network. In Proceedings of the 32nd ACM International Conference on Multimedia, Melbourne, Australia, pp. 4515\u20134523, 2024. DOI: https:\/\/doi.org\/10.1145\/3664647.3681088."},{"key":"1623_CR17","first-page":"6000","volume-title":"Proceedings of the 31st International Conference on Neural Information Processing Systems","author":"A Vaswani","year":"2017","unstructured":"A. Vaswani, N. Shazeer, N. Parmar, J. Uszkoreit, J. Jones, A. N. Gomez, \u0141. Kaiser, I. Polosukhin. Attention is all you need. In Proceedings of the 31st International Conference on Neural Information Processing Systems, Long Beach, USA, pp. 6000\u20136010, 2017."},{"issue":"10","key":"1623_CR18","doi-asserted-by":"publisher","first-page":"3365","DOI":"10.1109\/TPAMI.2020.2982166","volume":"43","author":"Z Wang","year":"2021","unstructured":"Z. Wang, J. Chen, S. C. H. Hoi. Deep learning for image super-resolution: A survey. IEEE Transactions on Pattern Analysis and Machine Intelligence, vol. 43, no. 10, pp. 3365\u20133387, 2021. DOI: https:\/\/doi.org\/10.1109\/TPAMI.2020.2982166.","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"key":"1623_CR19","doi-asserted-by":"publisher","unstructured":"J. Jiang, C. Wang, X. Liu, J. Ma. Deep learning-based face super-resolution: A survey. ACM Computing Surveys, vol. 55, no. 1, Article number 13, 2023. DOI: https:\/\/doi.org\/10.1145\/3485132.","DOI":"10.1145\/3485132"},{"key":"1623_CR20","volume-title":"Survey on deep face restoration: From non-blind to blind and beyond","author":"W Li","year":"2023","unstructured":"W. Li, M. Wang, K. Zhang, J. Li, X. Li, Y. Zhang, G. Gao, W. Deng, C. W. Lin. Survey on deep face restoration: From non-blind to blind and beyond, [Online], Available: https:\/\/arxiv.org\/abs\/2309.15490, 2023."},{"key":"1623_CR21","doi-asserted-by":"publisher","first-page":"4945","DOI":"10.1109\/CVPR46437.2021.00491","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"K C K Chan","year":"2021","unstructured":"K. C. K. Chan, X. Wang, K. Yu, C. Dong, C. C. Loy. BasicVSR: The search for essential components in video super-resolution and beyond. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, Nashville, USA, pp. 4945\u20134954, 2021. DOI: https:\/\/doi.org\/10.1109\/CVPR46437.2021.00491."},{"key":"1623_CR22","doi-asserted-by":"publisher","first-page":"5962","DOI":"10.1109\/CVPR52688.2022.00588","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"K C K Chan","year":"2022","unstructured":"K. C. K. Chan, S. Zhou, X. Xu, C. C. Loy. BasicVSR++: Improving video super-resolution with enhanced propagation and alignment. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, New Orleans, USA, pp. 5962\u20135971, 2022. DOI: https:\/\/doi.org\/10.1109\/CVPR52688.2022.00588."},{"key":"1623_CR23","doi-asserted-by":"publisher","first-page":"8934","DOI":"10.1109\/CVPR.2018.00931","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"D Sun","year":"2018","unstructured":"D. Sun, X. Yang, M. Y. Liu, J. Kautz. PWC-Net: CNNs for optical flow using pyramid, warping, and cost volume. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, Salt Lake City, USA, pp. 8934\u20138943, 2018. DOI: https:\/\/doi.org\/10.1109\/CVPR.2018.00931."},{"key":"1623_CR24","doi-asserted-by":"publisher","first-page":"1954","DOI":"10.1109\/CVPRW.2019.00247","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition Workshops","author":"X Wang","year":"2019","unstructured":"X. Wang, K. C. K. Chan, K. Yu, C. Dong, C. C. Loy. EDVR: Video restoration with enhanced deformable convolutional networks. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition Workshops, Long Beach, USA, pp. 1954\u20131963, 2019. DOI: https:\/\/doi.org\/10.1109\/CVPRW.2019.00247."},{"key":"1623_CR25","doi-asserted-by":"publisher","first-page":"3224","DOI":"10.1109\/CVPR.2018.00340","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"Y Jo","year":"2018","unstructured":"Y. Jo, S. W. Oh, J. Kang, S. J. Kim. Deep video superresolution network using dynamic upsampling filters without explicit motion compensation. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, Salt Lake City, USA, pp. 3224\u20133232, 2018. DOI: https:\/\/doi.org\/10.1109\/CVPR.2018.00340."},{"key":"1623_CR26","doi-asserted-by":"publisher","first-page":"2171","DOI":"10.1109\/TIP.2024.3372454","volume":"33","author":"J Liang","year":"2024","unstructured":"J. Liang, J. Cao, Y. Fan, K. Zhang, R. Ranjan, Y. Li, R. Timofte, L. Van Gool. VRT: A video restoration transformer. IEEE Transactions on Image Processing, vol. 33, pp. 2171\u20132182, 2024. DOI: https:\/\/doi.org\/10.1109\/TIP.2024.3372454.","journal-title":"IEEE Transactions on Image Processing"},{"issue":"9","key":"1623_CR27","doi-asserted-by":"publisher","first-page":"8929","DOI":"10.1109\/TCSVT.2025.3553160","volume":"35","author":"H Yue","year":"2025","unstructured":"H. Yue, C. Cao, L. Liao, J. Yang. RViDeformer: Efficient raw video denoising transformer with a larger benchmark dataset. IEEE Transactions on Circuits and Systems for Video Technology, vol. 35, no. 9, pp. 8929\u20138944, 2025. DOI: https:\/\/doi.org\/10.1109\/TCSVT.2025.3553160.","journal-title":"IEEE Transactions on Circuits and Systems for Video Technology"},{"issue":"8","key":"1623_CR28","doi-asserted-by":"publisher","first-page":"5981","DOI":"10.1007\/s10462-022-10147-y","volume":"55","author":"H Liu","year":"2022","unstructured":"H. Liu, Z. Ruan, P. Zhao, C. Dong, F. Shang, Y. Liu, L. Yang, R. Timofte. Video super-resolution based on deep learning: A comprehensive survey. Artificial Intelligence Review, vol. 55, no. 8, pp. 5981\u20136035, 2022. DOI: https:\/\/doi.org\/10.1007\/s10462-022-10147-y.","journal-title":"Artificial Intelligence Review"},{"issue":"4","key":"1623_CR29","doi-asserted-by":"publisher","first-page":"2655","DOI":"10.1109\/TETCI.2024.3398015","volume":"8","author":"A A Baniya","year":"2024","unstructured":"A. A. Baniya, T. K. Lee, P. W. Eklund, S. Aryal. A survey of deep learning video super-resolution. IEEE Transactions on Emerging Topics in Computational Intelligence, vol. 8, no. 4, pp. 2655\u20132676, 2024. DOI: https:\/\/doi.org\/10.1109\/TETCI.2024.3398015.","journal-title":"IEEE Transactions on Emerging Topics in Computational Intelligence"},{"key":"1623_CR30","doi-asserted-by":"publisher","first-page":"3","DOI":"10.1007\/978-3-031-27066-6_1","volume-title":"Proceedings of the 16th Asian Conference on Computer Vision","author":"S Bian","year":"2022","unstructured":"S. Bian, H. Li, F. Yu, J. Liu, C. Song, Y. Tang. FAPN: Face alignment propagation network for face video super-resolution. In Proceedings of the 16th Asian Conference on Computer Vision, Macao, China, pp. 3\u201318, 2022. DOI: https:\/\/doi.org\/10.1007\/978-3-031-27066-6_1."},{"key":"1623_CR31","doi-asserted-by":"publisher","DOI":"10.1109\/IJCNN60899.2024.10650940","volume-title":"Proceedings of International Joint Conference on Neural Networks","author":"R Shi","year":"2024","unstructured":"R. Shi, W. Guo, S. Ge. Unsupervised video face superresolution via untrained neural network priors. In Proceedings of International Joint Conference on Neural Networks, Yokohama, Japan, 2024. DOI: https:\/\/doi.org\/10.1109\/IJCNN60899.2024.10650940."},{"key":"1623_CR32","doi-asserted-by":"publisher","first-page":"5228","DOI":"10.1109\/WACV61041.2025.00511","volume-title":"Proceedings of IEEE\/CVF Winter Conference on Applications of Computer Vision","author":"Z Zou","year":"2025","unstructured":"Z. Zou, J. Liu, S. Shoushtari, Y. Wang, U. S. Kamilov. FLAIR: A conditional diffusion framework with applications to face video restoration. In Proceedings of IEEE\/CVF Winter Conference on Applications of Computer Vision, Tucson, USA, pp. 5228\u20135238, 2025. DOI: https:\/\/doi.org\/10.1109\/WACV61041.2025.00511."},{"key":"1623_CR33","first-page":"367","volume-title":"Proceedings of the 16th Asian Conference on Machine Learning","author":"Y Xu","year":"2025","unstructured":"Y. Xu, Z. Song, J. Lu. Universal video face restoration method based on vision-language model. In Proceedings of the 16th Asian Conference on Machine Learning, Hanoi, Vietnam, pp. 367\u2013382, 2025."},{"key":"1623_CR34","doi-asserted-by":"publisher","first-page":"2183","DOI":"10.1109\/CVPR52734.2025.00209","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"Y Wang","year":"2025","unstructured":"Y. Wang, J. Teng, J. Cao, Y. Li, C. Ma, H. Xu, D. Luo. Efficient video face enhancement with enhanced spatial-temporal consistency. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, Nashville, USA, pp. 2183\u20132193, 2025. DOI: https:\/\/doi.org\/10.1109\/CVPR52734.2025.00209."},{"key":"1623_CR35","doi-asserted-by":"publisher","first-page":"1417","DOI":"10.1145\/3664647.3680917","volume-title":"Proceedings of the 32nd ACM International Conference on Multimedia","author":"J Tan","year":"2024","unstructured":"J. Tan, H. Park, Y. Zhang, T. Wang, K. Zhang, X. Kong, P. Dai, Z. Liu, W. Luo. Blind face video restoration with temporal consistent generative prior and degradation-aware prompt. In Proceedings of the 32nd ACM International Conference on Multimedia, Melbourne, Australia, pp. 1417\u20131426, 2024. DOI: https:\/\/doi.org\/10.1145\/3664647.3680917."},{"key":"1623_CR36","doi-asserted-by":"publisher","first-page":"7406","DOI":"10.1109\/CVPR52734.2025.00694","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"Z Wang","year":"2025","unstructured":"Z. Wang, X. Chen, C. Xu, J. Zhu, X. Hu, J. Zhang, C. Wang, Y. Liu, Y. Zhou, R. Ji. SVFR: A unified framework for generalized video face restoration. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, Nashville, USA, pp. 7406\u20137415, 2025. DOI: https:\/\/doi.org\/10.1109\/CVPR52734.2025.00694."},{"key":"1623_CR37","doi-asserted-by":"publisher","DOI":"10.1609\/aaai.v39i8.32944","volume-title":"Proceedings of the 39th AAAI Conference on Artificial Intelligence","author":"L Xie","year":"2025","unstructured":"L. Xie, B. Zheng, W. Xue, Y. Zhang, L. Jiang, R. Xu, S. Wu, H. S. Wong. Discrete prior-based temporal-coherent content prediction for blind face video restoration. In Proceedings of the 39th AAAI Conference on Artificial Intelligence, Philadelphia, USA, Article number 971, 2025. DOI: https:\/\/doi.org\/10.1609\/aaai.v39i8.32944."},{"key":"1623_CR38","doi-asserted-by":"publisher","first-page":"17821","DOI":"10.1109\/CVPR52734.2025.01660","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"L Xie","year":"2025","unstructured":"L. Xie, B. Zheng, S. Wu, H. S. Wong. Dynamic content prediction with motion-aware priors for blind face video restoration. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, Nashville, USA, pp. 17821\u201317830, 2025. DOI: https:\/\/doi.org\/10.1109\/CVPR52734.2025.01660."},{"key":"1623_CR39","first-page":"1489","volume-title":"Proceedings of the 33rd International Joint Conference on Artificial Intelligence","author":"K Xu","year":"2024","unstructured":"K. Xu, L. Xu, G. He, W. Yu, Y. Li. Beyond alignment: Blind video face restoration via parsing-guided temporal-coherent transformer. In Proceedings of the 33rd International Joint Conference on Artificial Intelligence, Jeju, Republic of Korea, pp. 1489\u20131497, 2024."},{"key":"1623_CR40","doi-asserted-by":"publisher","first-page":"202","DOI":"10.1007\/978-3-031-73347-5_12","volume-title":"Proceedings of the 18th European Conference on Computer Vision","author":"R Feng","year":"2024","unstructured":"R. Feng, C. Li, C. C. Loy. Kalman-inspired feature propagation for video face super-resolution. In Proceedings of the 18th European Conference on Computer Vision, Milan, Italy, pp. 202\u2013218, 2024. DOI: https:\/\/doi.org\/10.1007\/978-3-031-73347-5_12."},{"key":"1623_CR41","first-page":"2616","volume-title":"Proceedings of the 18th Annual Conference of the International Speech Communication Association","author":"A Nagrani","year":"2017","unstructured":"A. Nagrani, J S. Chung, A. Zisserman. VoxCeleb: A large-scale speaker identification dataset. In Proceedings of the 18th Annual Conference of the International Speech Communication Association, Stockholm, Sweden, pp. 2616\u20132620, 2017."},{"key":"1623_CR42","first-page":"1086","volume-title":"Proceedings of the 19th Annual Conference of the International Speech Communication Association","author":"J S Chung","year":"2018","unstructured":"J. S. Chung, A. Nagrani, A. Zisserman. VoxCeleb2: Deep speaker recognition. In Proceedings of the 19th Annual Conference of the International Speech Communication Association, Hyderabad, India, pp. 1086\u20131090, 2018."},{"key":"1623_CR43","doi-asserted-by":"publisher","first-page":"650","DOI":"10.1007\/978-3-031-20071-7_38","volume-title":"Proceedings of the 17th European Conference on Computer Vision","author":"H Zhu","year":"2022","unstructured":"H. Zhu, W. Wu, W. Zhu, L. Jiang, S. Tang, L. Zhang, Z. Liu, C C. Loy. CelebV-HQ: A large-scale video facial attributes dataset. In Proceedings of the 17th European Conference on Computer Vision, Tel Aviv, Israel, pp. 650\u2013667, 2022. DOI: https:\/\/doi.org\/10.1007\/978-3-031-20071-7_38."},{"key":"1623_CR44","doi-asserted-by":"publisher","first-page":"14805","DOI":"10.1109\/CVPR52729.2023.01422","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"J Yu","year":"2023","unstructured":"J. Yu, H. Zhu, L. Jiang, C. C. Loy, W. Cai, W. Wu. CelebV-Text: A large-scale facial text-video dataset. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, Vancouver, Canada, pp. 14805\u201314814, 2023. DOI: https:\/\/doi.org\/10.1109\/CVPR52729.2023.01422."},{"key":"1623_CR45","doi-asserted-by":"publisher","first-page":"656","DOI":"10.1109\/CVPRW56347.2022.00081","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition Workshops","author":"L Xie","year":"2022","unstructured":"L. Xie, X. Wang, H. Zhang, C. Dong, Y. Shan. VFHQ: A high-quality dataset and benchmark for video face super-resolution. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition Workshops, New Orleans, USA, pp. 656\u2013665, 2022. DOI: https:\/\/doi.org\/10.1109\/CVPRW56347.2022.00081."},{"key":"1623_CR46","doi-asserted-by":"publisher","first-page":"1003","DOI":"10.1109\/ICCVW.2015.132","volume-title":"Proceedings of IEEE International Conference on Computer Vision Workshop","author":"J Shen","year":"2015","unstructured":"J. Shen, S. Zafeiriou, G. G. Chrysos, J. Kossaifi, G. Tzimiropoulos, M. Pantic. The first facial landmark tracking in-the-wild challenge: Benchmark and results. In Proceedings of IEEE International Conference on Computer Vision Workshop, Santiago, Chile, pp. 1003\u20131011, 2015. DOI: https:\/\/doi.org\/10.1109\/ICCVW.2015.132."},{"key":"1623_CR47","doi-asserted-by":"publisher","unstructured":"S. Suwajanakorn, S. M. Seitz, I. Kemelmacher-Shlizerman. Synthesizing Obama: Learning lip sync from audio. ACM Transactions on Graphics, vol. 36, no. 4, Article number 95, 2017. DOI: https:\/\/doi.org\/10.1145\/3072959.3073640.","DOI":"10.1145\/3072959.3073640"},{"key":"1623_CR48","doi-asserted-by":"publisher","first-page":"397","DOI":"10.1109\/ICCVW.2013.59","volume-title":"Proceedings of IEEE International Conference on Computer Vision Workshops","author":"C Sagonas","year":"2013","unstructured":"C. Sagonas, G. Tzimiropoulos, S. Zafeiriou, M. Pantic. 300 faces in-the-wild challenge: The first facial landmark localization challenge. In Proceedings of IEEE International Conference on Computer Vision Workshops, Sydney, Australia, pp. 397\u2013403, 2013. DOI: https:\/\/doi.org\/10.1109\/ICCVW.2013.59."},{"issue":"3","key":"1623_CR49","doi-asserted-by":"publisher","first-page":"81","DOI":"10.1109\/97.995823","volume":"9","author":"Z Wang","year":"2002","unstructured":"Z. Wang, A. C. Bovik. A universal image quality index. IEEE Signal Processing Letters, vol. 9, no. 3, pp. 81\u201384, 2002. DOI: https:\/\/doi.org\/10.1109\/97.995823.","journal-title":"IEEE Signal Processing Letters"},{"key":"1623_CR50","doi-asserted-by":"publisher","first-page":"586","DOI":"10.1109\/CVPR.2018.00068","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"R Zhang","year":"2018","unstructured":"R. Zhang, P. Isola, A. A. Efros, E. Shechtman, O. Wang. The unreasonable effectiveness of deep features as a perceptual metric. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, Salt Lake City, USA, pp. 586\u2013595, 2018. DOI: https:\/\/doi.org\/10.1109\/CVPR.2018.00068."},{"key":"1623_CR51","first-page":"6629","volume-title":"Proceedings of the 31st International Conference on Neural Information Processing Systems","author":"M Heusel","year":"2017","unstructured":"M. Heusel, H. Ramsauer, T. Unterthiner, B. Nessler, S. Hochreiter. GANs trained by a two time-scale update rule converge to a local Nash equilibrium. In Proceedings of the 31st International Conference on Neural Information Processing Systems, Long Beach, USA, pp. 6629\u20136640, 2017."},{"key":"1623_CR52","doi-asserted-by":"publisher","first-page":"2818","DOI":"10.1109\/CVPR.2016.308","volume-title":"Proceedings of IEEE Conference on Computer Vision and Pattern Recognition","author":"C Szegedy","year":"2016","unstructured":"C. Szegedy, V. Vanhoucke, S. Ioffe, J. Shlens, Z. Wojna. Rethinking the inception architecture for computer vision. In Proceedings of IEEE Conference on Computer Vision and Pattern Recognition, Las Vegas, USA, pp. 2818\u20132826, 2016. DOI: https:\/\/doi.org\/10.1109\/CVPR.2016.308."},{"key":"1623_CR53","volume-title":"Proceedings of the 33rd International Conference on Neural Information Processing Systems","author":"A Siarohin","year":"2019","unstructured":"A. Siarohin, S. Lathuili\u00e8re, S. Tulyakov, E. Ricci, N. Sebe. First order motion model for image animation. In Proceedings of the 33rd International Conference on Neural Information Processing Systems, Vancouver, Canada, Article number 641, 2019."},{"key":"1623_CR54","doi-asserted-by":"publisher","first-page":"6970","DOI":"10.1109\/ICCV.2019.00707","volume-title":"Proceedings of IEEE\/CVF International Conference on Computer Vision","author":"X Wang","year":"2019","unstructured":"X. Wang, L. Bo, F. Li. Adaptive wing loss for robust face alignment via heatmap regression. In Proceedings of IEEE\/CVF International Conference on Computer Vision, Seoul, Republic of Korea, pp. 6970\u20136980, 2019. DOI: https:\/\/doi.org\/10.1109\/ICCV.2019.00707."},{"key":"1623_CR55","doi-asserted-by":"publisher","first-page":"4685","DOI":"10.1109\/CVPR.2019.00482","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"J Deng","year":"2019","unstructured":"J. Deng, J. Guo, N. Xue, S. Zafeiriou. ArcFace: Additive angular margin loss for deep face recognition. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, Long Beach, USA, pp. 4685\u20134694, 2019. DOI: https:\/\/doi.org\/10.1109\/CVPR.2019.00482."},{"key":"1623_CR56","volume-title":"Towards accurate generative models of video: A new metric & challenges","author":"T Unterthiner","year":"2018","unstructured":"T. Unterthiner, S. Van Steenkiste, K. Kurach, R. Marinier, M. Michalski, S. Gelly. Towards accurate generative models of video: A new metric & challenges, [Online], Available: https:\/\/arxiv.org\/abs\/1812.01717, 2018."},{"key":"1623_CR57","doi-asserted-by":"publisher","first-page":"4724","DOI":"10.1109\/CVPR.2017.502","volume-title":"Proceedings of IEEE Conference on Computer Vision and Pattern Recognition","author":"J Carreira","year":"2017","unstructured":"J. Carreira, A. Zisserman. Quo Vadis, action recognition? A new model and the kinetics dataset. In Proceedings of IEEE Conference on Computer Vision and Pattern Recognition, Honolulu, USA, pp. 4724\u20134733, 2017. DOI: https:\/\/doi.org\/10.1109\/CVPR.2017.502."},{"key":"1623_CR58","doi-asserted-by":"publisher","first-page":"179","DOI":"10.1007\/978-3-030-01267-0_11","volume-title":"Proceedings of the 15th European Conference on Computer Vision","author":"W S Lai","year":"2018","unstructured":"W. S. Lai, J. B. Huang, O. Wang, E. Shechtman, E. Yumer, M. H. Yang. Learning blind video temporal consistency. In Proceedings of the 15th European Conference on Computer Vision, Munich, Germany, pp. 179\u2013195, 2018. DOI: https:\/\/doi.org\/10.1007\/978-3-030-01267-0_11."},{"key":"1623_CR59","doi-asserted-by":"publisher","first-page":"22139","DOI":"10.1109\/CVPR52733.2024.02090","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"Y Liu","year":"2024","unstructured":"Y. Liu, X. Cun, X. Liu, X. Wang, Y. Zhang, H. Chen, Y. Liu, T. Zeng, R. Chan, Y. Shan. EvalCrafter: Benchmarking and evaluating large video generation models. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, Seattle, USA, pp. 22139\u201322149, 2024. DOI: https:\/\/doi.org\/10.1109\/CVPR52733.2024.02090."},{"key":"1623_CR60","doi-asserted-by":"publisher","first-page":"402","DOI":"10.1007\/978-3-030-58536-5_24","volume-title":"Proceedings of the 16th European Conference on Computer Vision","author":"Z Teed","year":"2020","unstructured":"Z. Teed, J. Deng. RAFT: Recurrent all-pairs field transforms for optical flow. In Proceedings of the 16th European Conference on Computer Vision, Glasgow, UK, pp. 402\u2013419, 2020. DOI: https:\/\/doi.org\/10.1007\/978-3-030-58536-5_24."},{"issue":"1","key":"1623_CR61","first-page":"301","volume":"2","author":"E Donelly","year":"2012","unstructured":"E. Donelly. Very deep convolutional networks for large-scale image recognition. International Journal of Artificial Intelligence and Machine Learning, vol. 2, no. 1, pp. 301\u2013307, 2012.","journal-title":"International Journal of Artificial Intelligence and Machine Learning"},{"key":"1623_CR62","doi-asserted-by":"publisher","first-page":"9387","DOI":"10.1109\/ICCV.2019.00948","volume-title":"Proceedings of IEEE\/CVF International Conference on Computer Vision","author":"W Ren","year":"2019","unstructured":"W. Ren, J. Yang, S. Deng, D. Wipf, X. Cao, X. Tong. Face video deblurring using 3D facial priors. In Proceedings of IEEE\/CVF International Conference on Computer Vision, Seoul, Republic of Korea, pp. 9387\u20139396, 2019. DOI: https:\/\/doi.org\/10.1109\/ICCV.2019.00948."},{"key":"1623_CR63","doi-asserted-by":"publisher","first-page":"357","DOI":"10.1007\/978-981-10-7302-1_30","volume-title":"Proceedings of the 2nd CCF Chinese Conference on Computer Vision","author":"D Li","year":"2017","unstructured":"D. Li, Z. Wang. Face video super-resolution with identity guided generative adversarial networks. In Proceedings of the 2nd CCF Chinese Conference on Computer Vision, Tianjin, China, pp. 357\u2013369, 2017. DOI: https:\/\/doi.org\/10.1007\/978-981-10-7302-1_30."},{"key":"1623_CR64","doi-asserted-by":"publisher","first-page":"17601","DOI":"10.1109\/CVPR52688.2022.01710","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition Workshops","author":"P Dai","year":"2022","unstructured":"P. Dai, X. Yu, L. Ma, B. Zhang, J. Li, W. Li, J. Shen, X. Qi. Video demoir\u00e9ing with relation-based temporal consistency. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition Workshops, New Orleans, USA, pp. 17601\u201317610, 2022. DOI: https:\/\/doi.org\/10.1109\/CVPR52688.2022.01710."},{"key":"1623_CR65","doi-asserted-by":"publisher","first-page":"1513","DOI":"10.1109\/ICCVW54120.2021.00176","volume-title":"Proceedings of IEEE\/CVF International Conference on Computer Vision Workshops","author":"F Yu","year":"2021","unstructured":"F. Yu, H. Li, S. Bian, Y. Tang. An efficient network design for face video super-resolution. In Proceedings of IEEE\/CVF International Conference on Computer Vision Workshops, Montreal, Canada, pp. 1513\u20131520, 2021. DOI: https:\/\/doi.org\/10.1109\/ICCVW54120.2021.00176."},{"key":"1623_CR66","doi-asserted-by":"publisher","first-page":"12468","DOI":"10.1609\/aaai.v34i07.6934","volume-title":"Proceedings of the 34th AAAI Conference on Artificial Intelligence","author":"J Xin","year":"2020","unstructured":"J. Xin, N. Wang, J. Li, X. Gao, Z. Li. Video face superresolution with motion-adaptive feedback cell. In Proceedings of the 34th AAAI Conference on Artificial Intelligence, New York, USA, pp. 12468\u201312475, 2020. DOI: https:\/\/doi.org\/10.1609\/aaai.v34i07.6934."},{"issue":"2","key":"1623_CR67","doi-asserted-by":"publisher","first-page":"2024","DOI":"10.1109\/TPAMI.2022.3157388","volume":"45","author":"X Zhang","year":"2023","unstructured":"X. Zhang, X. Wu. Multi-modality deep restoration of extremely compressed face videos. IEEE Transactions on Pattern Analysis and Machine Intelligence, vol. 45, no. 2, pp. 2024\u20132037, 2023. DOI: https:\/\/doi.org\/10.1109\/TPAMI.2022.3157388.","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"key":"1623_CR68","doi-asserted-by":"publisher","first-page":"5676","DOI":"10.1109\/TIP.2024.3463414","volume":"33","author":"Z Wang","year":"2024","unstructured":"Z. Wang, J. Zhang, X. Wang, T. Chen, Y. Shan, W. Wang, P. Luo. Analysis and benchmarking of extending blind face image restoration to videos. IEEE Transactions on Image Processing, vol. 33, pp. 5676\u20135687, 2024. DOI: https:\/\/doi.org\/10.1109\/TIP.2024.3463414.","journal-title":"IEEE Transactions on Image Processing"},{"key":"1623_CR69","doi-asserted-by":"publisher","DOI":"10.1109\/ICCUBEA47591.2019.9128399","volume-title":"Proceedings of the 5th International Conference on Computing, Communication, Control and Automation","author":"A B Deshmukh","year":"2019","unstructured":"A. B. Deshmukh, N. U. Rani. Face video super resolution using deep convolutional neural network. In Proceedings of the 5th International Conference on Computing, Communication, Control and Automation, Pune, India, 2019. DOI: https:\/\/doi.org\/10.1109\/ICCUBEA47591.2019.9128399."},{"key":"1623_CR70","doi-asserted-by":"publisher","first-page":"3078","DOI":"10.1109\/TIP.2019.2955640","volume":"29","author":"C Fang","year":"2020","unstructured":"C. Fang, G. Li, X. Han, Y. Yu. Self-enhanced convolutional network for facial video hallucination. IEEE Transactions on Image Processing, vol. 29, pp. 3078\u20133090, 2020. DOI: https:\/\/doi.org\/10.1109\/TIP.2019.2955640.","journal-title":"IEEE Transactions on Image Processing"},{"key":"1623_CR71","doi-asserted-by":"publisher","first-page":"12868","DOI":"10.1109\/CVPR46437.2021.01268","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"P Esser","year":"2021","unstructured":"P. Esser, R. Rombach, B. Ommer. Taming transformers for high-resolution image synthesis. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, Nashville, USA, pp. 12868\u201312878, 2021. DOI: https:\/\/doi.org\/10.1109\/CVPR46437.2021.01268."},{"key":"1623_CR72","doi-asserted-by":"publisher","first-page":"205","DOI":"10.1007\/978-3-031-25066-8_9","volume-title":"Proceedings of European Conference on Computer Vision","author":"H Cao","year":"2022","unstructured":"H. Cao, Y. Wang, J. Chen, D. Jiang, X. Zhang, Q. Tian, M. Wang. Swin-Unet: Unet-like pure transformer for medical image segmentation. In Proceedings of European Conference on Computer Vision, Tel Aviv, Israel, pp. 205\u2013218, 2022. DOI: https:\/\/doi.org\/10.1007\/978-3-031-25066-8_9."},{"key":"1623_CR73","first-page":"6309","volume-title":"Proceedings of the 31st International Conference on Neural Information Processing Systems","author":"A van den Oord","year":"2017","unstructured":"A. van den Oord, O. Vinyals, K. Kavukcuoglu. Neural discrete representation learning. In Proceedings of the 31st International Conference on Neural Information Processing Systems, Long Beach, USA, pp. 6309\u20136318, 2017."},{"key":"1623_CR74","first-page":"14507","volume-title":"Proceedings of IEEE\/CVF International Conference on Computer Vision","author":"B Chen","year":"2025","unstructured":"B. Chen, C. Liu, W. Yuan, Z. Dong, S. Zhu. Dirichlet-constrained variational codebook learning for temporally coherent video face restoration. In Proceedings of IEEE\/CVF International Conference on Computer Vision, Honolulu, USA, pp. 14507\u201314516, 2025."},{"key":"1623_CR75","doi-asserted-by":"publisher","first-page":"10674","DOI":"10.1109\/CVPR52688.2022.01042","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"R Rombach","year":"2022","unstructured":"R. Rombach, A. Blattmann, D. Lorenz, P. Esser, B. Ommer. High-resolution image synthesis with latent diffusion models. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, New Orleans, USA, pp. 10674\u201310685, 2022. DOI: https:\/\/doi.org\/10.1109\/CVPR52688.2022.01042."},{"key":"1623_CR76","doi-asserted-by":"publisher","first-page":"6108","DOI":"10.1109\/CVPRW63382.2024.00617","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition Workshops","author":"Z Chen","year":"2024","unstructured":"Z. Chen, Z. Wu, E. Zamfir, K. Zhang, Y. Zhang, R. Timofte. NTIRE 2024 challenge on image super-resolution (\u00d74): Methods and results. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition Workshops, Seattle, USA, pp. 6108\u20136132, 2024. DOI: https:\/\/doi.org\/10.1109\/CVPRW63382.2024.00617."},{"key":"1623_CR77","doi-asserted-by":"publisher","first-page":"9232","DOI":"10.1109\/CVPR52733.2024.00882","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"Z Chen","year":"2024","unstructured":"Z. Chen, F. Long, Z. Qiu, T. Yao, W. Zhou, J. Luo, T. Mei. Learning spatial adaptation and temporal coherence in diffusion models for video super-resolution. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, Seattle, USA, pp. 9232\u20139241, 2024. DOI: https:\/\/doi.org\/10.1109\/CVPR52733.2024.00882."},{"key":"1623_CR78","doi-asserted-by":"publisher","unstructured":"X. Zhang, J. Ma, G. Wang, Q. Zhang, H. Zhang, L. Zhang. Perceive-IR: Learning to perceive degradation better for all-in-one image restoration. IEEE Transactions on Image Processing, 2025. DOI: https:\/\/doi.org\/10.1109\/TIP.2025.3566300.","DOI":"10.1109\/TIP.2025.3566300"},{"key":"1623_CR79","volume-title":"Proceedings of the 34th International Conference on Neural Information Processing Systems","author":"J Ho","year":"2020","unstructured":"J. Ho, A. Jain, P. Abbeel. Denoising diffusion probabilistic models. In Proceedings of the 34th International Conference on Neural Information Processing Systems, Vancouver, Canada, Article number 574, 2020."},{"key":"1623_CR80","volume-title":"Proceedings of the 35th International Conference on Neural Information Processing Systems","author":"P Dhariwal","year":"2021","unstructured":"P. Dhariwal, A. Nichol. Diffusion models beat GANs on image synthesis. In Proceedings of the 35th International Conference on Neural Information Processing Systems, Article number 672, 2021."},{"key":"1623_CR81","doi-asserted-by":"publisher","unstructured":"S. Cai, X. Ding, J. Fang, Q. Liu, Y. Yang. Spatially adaptive representation of facial meshes for face video super-resolution. Expert Systems with Applications, vol. 296, Article number 128864, 2026. DOI: https:\/\/doi.org\/10.1016\/j.eswa.2025.128864.","DOI":"10.1016\/j.eswa.2025.128864"},{"issue":"1","key":"1623_CR82","doi-asserted-by":"publisher","first-page":"83","DOI":"10.26599\/CVM.2025.9450383","volume":"11","author":"Y Q Yang","year":"2025","unstructured":"Y. Q. Yang, Y. X. Guo, J. Y. Xiong, Y. Liu, H. Pan, P. S. Wang, X. Tong, B. Guo. Swin3D: A pretrained transformer backbone for 3D indoor scene understanding. Computational Visual Media, vol. 11, no. 1, pp. 83\u2013101, 2025. DOI: https:\/\/doi.org\/10.26599\/CVM.2025.9450383.","journal-title":"Computational Visual Media"},{"key":"1623_CR83","doi-asserted-by":"publisher","first-page":"234","DOI":"10.1007\/978-3-319-24574-4_28","volume-title":"Proceedings of the 18th International Conference on Medical Image Computing and Computer-Assisted Intervention","author":"O Ronneberger","year":"2015","unstructured":"O. Ronneberger, P. Fischer, T. Brox. U-Net: Convolutional networks for biomedical image segmentation. In Proceedings of the 18th International Conference on Medical Image Computing and Computer-Assisted Intervention, Munich, Germany, pp. 234\u2013241, 2015. DOI: https:\/\/doi.org\/10.1007\/978-3-319-24574-4_28."},{"key":"1623_CR84","volume-title":"Stable video diffusion: Scaling latent video diffusion models to large datasets","author":"A Blattmann","year":"2023","unstructured":"A. Blattmann, T. Dockhorn, S. Kulal, D. Mendelevitch, M. Kilian, D. Lorenz, Y. Levi, Z. English, V. Voleti, A. Letts, V. Jampani, R. Rombach. Stable video diffusion: Scaling latent video diffusion models to large datasets, [Online], Available: https:\/\/arxiv.org\/abs\/2311.15127, 2023."},{"issue":"12","key":"1623_CR85","doi-asserted-by":"publisher","first-page":"15462","DOI":"10.1109\/TPAMI.2023.3315753","volume":"45","author":"Z Wang","year":"2023","unstructured":"Z. Wang, J. Zhang, T. Chen, W. Wang, P. Luo. Restore-Former++: Towards real-world blind face restoration from undegraded key-value pairs. IEEE Transactions on Pattern Analysis and Machine Intelligence, vol. 45, no. 12, pp. 15462\u201315476, 2023. DOI: https:\/\/doi.org\/10.1109\/TPAMI.2023.3315753.","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"key":"1623_CR86","doi-asserted-by":"publisher","first-page":"5952","DOI":"10.1109\/CVPR52688.2022.00587","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"K C K Chan","year":"2022","unstructured":"K. C. K. Chan, S. Zhou, X. Xu, C. C. Loy. Investigating tradeoffs in real-world video super-resolution. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, New Orleans, USA, pp. 5952\u20135961, 2022. DOI: https:\/\/doi.org\/10.1109\/CVPR52688.2022.00587."},{"key":"1623_CR87","doi-asserted-by":"publisher","first-page":"9059","DOI":"10.1109\/CVPR52733.2024.00865","volume-title":"Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition","author":"Z Chaouai","year":"2024","unstructured":"Z. Chaouai, M. Tamaazousti. Universal robustness via median randomized smoothing for real-world super-resolution. In Proceedings of IEEE\/CVF Conference on Computer Vision and Pattern Recognition, Seattle, USA, pp. 9059\u20139068, 2024. DOI: https:\/\/doi.org\/10.1109\/CVPR52733.2024.00865."},{"key":"1623_CR88","volume-title":"Can no-reference quality-assessment methods serve as perceptual losses for super-resolution?","author":"E Kashkarov","year":"2024","unstructured":"E. Kashkarov, E. Chistov, I. Molodetskikh, D. Vatolin. Can no-reference quality-assessment methods serve as perceptual losses for super-resolution? [Online], Available: https:\/\/arxiv.org\/abs\/2405.20392, 2024."},{"key":"1623_CR89","volume-title":"Proceedings of the 13th International Conference on Learning Representations","author":"K Tao","year":"2025","unstructured":"K. Tao, J. Gu, Y. Zhang, X. Wang, N. Cheng. Overcoming false illusions in real-world face restoration with multi-modal guided diffusion model. In Proceedings of the 13th International Conference on Learning Representations, Singapore, 2025."},{"key":"1623_CR90","doi-asserted-by":"publisher","first-page":"2564","DOI":"10.1145\/3664647.3680603","volume-title":"Proceedings of the 32nd ACM International Conference on Multimedia","author":"Z Yin","year":"2024","unstructured":"Z. Yin, M. Ma, G. Lin, Y. Zheng. Exploring data efficiency in image restoration: A Gaussian denoising case study. In Proceedings of the 32nd ACM International Conference on Multimedia, Melbourne, Australia, pp. 2564\u20132573, 2024. DOI: https:\/\/doi.org\/10.1145\/3664647.3680603."}],"container-title":["Machine Intelligence Research"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11633-025-1623-x.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11633-025-1623-x","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11633-025-1623-x.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2026,6,2]],"date-time":"2026-06-02T15:03:53Z","timestamp":1780412633000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11633-025-1623-x"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,5,11]]},"references-count":90,"journal-issue":{"issue":"3","published-print":{"date-parts":[[2026,6]]}},"alternative-id":["1623"],"URL":"https:\/\/doi.org\/10.1007\/s11633-025-1623-x","relation":{},"ISSN":["2731-538X","2731-5398"],"issn-type":[{"value":"2731-538X","type":"print"},{"value":"2731-5398","type":"electronic"}],"subject":[],"published":{"date-parts":[[2026,5,11]]},"assertion":[{"value":"5 July 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"25 November 2025","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"11 May 2026","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"The authors declared that they have no conflicts of interest to this work.","order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations of conflict of interest"}}]}}