{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2025,11,5]],"date-time":"2025-11-05T18:45:10Z","timestamp":1762368310421,"version":"build-2065373602"},"reference-count":46,"publisher":"Tsinghua University Press","issue":"6","license":[{"start":{"date-parts":[[2024,12,1]],"date-time":"2024-12-01T00:00:00Z","timestamp":1733011200000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/creativecommons.org\/licenses\/by\/4.0"},{"start":{"date-parts":[[2024,3,22]],"date-time":"2024-03-22T00:00:00Z","timestamp":1711065600000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/creativecommons.org\/licenses\/by\/4.0"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Comp. Visual. Med."],"published-print":{"date-parts":[[2024,12]]},"DOI":"10.1007\/s41095-023-0379-8","type":"journal-article","created":{"date-parts":[[2024,3,22]],"date-time":"2024-03-22T04:16:22Z","timestamp":1711080982000},"page":"1063-1078","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":2,"title":["Geometry-aware 3D pose transfer using transformer autoencoder"],"prefix":"10.26599","volume":"10","author":[{"given":"Shanghuan","family":"Liu","sequence":"first","affiliation":[{"name":"Key Laboratory of Measurement and Control of Complex Systems of Engineering, Ministry of Education, Southeast University, Nanjing, 210096, China"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Shaoyan","family":"Gai","sequence":"additional","affiliation":[{"name":"Key Laboratory of Measurement and Control of Complex Systems of Engineering, Ministry of Education, Southeast University, Nanjing, 210096, China"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Feipeng","family":"Da","sequence":"additional","affiliation":[{"name":"Key Laboratory of Measurement and Control of Complex Systems of Engineering, Ministry of Education, Southeast University, Nanjing, 210096, China"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Fazal","family":"Waris","sequence":"additional","affiliation":[{"name":"Key Laboratory of Measurement and Control of Complex Systems of Engineering, Ministry of Education, Southeast University, Nanjing, 210096, China"}],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"11138","reference":[{"key":"379_CR1","doi-asserted-by":"publisher","first-page":"46","DOI":"10.1016\/j.cag.2022.03.007","volume":"104","author":"Y P Ye","year":"2022","unstructured":"Ye, Y. P.; Song, Z.; Zhao, J. High-fidelity 3D real-time facial animation using infrared structured light sensing system. Computers & Graphics Vol. 104, 46\u201358, 2022.","journal-title":"Computers & Graphics"},{"key":"379_CR2","doi-asserted-by":"publisher","first-page":"52","DOI":"10.1016\/j.cag.2020.10.004","volume":"94","author":"R A Roberts","year":"2021","unstructured":"Roberts, R. A.; dos Anjos, R. K.; Maejima, A.; Anjyo, K. Deformation transfer survey. Computers & Graphics Vol. 94, 52\u201361, 2021.","journal-title":"Computers & Graphics"},{"key":"379_CR3","doi-asserted-by":"crossref","unstructured":"Ben-Chen, M.; Weber, O.; Gotsman, C. Spatial deformation transfer. In: Proceedings of the ACM SIGGRAPH\/Eurographics Symposium on Computer Animation, 67\u201374, 2009.","DOI":"10.1145\/1599470.1599479"},{"issue":"2","key":"379_CR4","first-page":"379","volume":"26","author":"H K Chu","year":"2010","unstructured":"Chu, H. K.; Lin, C. H. Example-based deformation transfer for 3D polygon models. Journal of Information Science and Engineering Vol. 26, No. 2, 379\u2013391, 2010.","journal-title":"Journal of Information Science and Engineering"},{"key":"379_CR5","doi-asserted-by":"publisher","first-page":"167","DOI":"10.1016\/j.cag.2020.05.013","volume":"89","author":"Y Z Zhang","year":"2020","unstructured":"Zhang, Y. Z.; Zheng, J. M.; Cai, Y. Y. Proxy-driven free-form deformation by topology-adjustable control lattice. Computers & Graphics Vol. 89, 167\u2013177, 2020.","journal-title":"Computers & Graphics"},{"key":"379_CR6","series-title":"Lecture Notes in Computer Science","doi-asserted-by":"publisher","first-page":"640","DOI":"10.1007\/978-3-031-20086-1_37","volume-title":"Computer Vision\u2013ECCV 2022","author":"Z Liao","year":"2022","unstructured":"Liao, Z.; Yang, J. M.; Saito, J.; Pons-Moll, G.; Zhou, Y. Skeleton-free pose transfer for stylized 3D characters. In: Computer Vision\u2013ECCV 2022. Lecture Notes in Computer Science, Vol. 13662. Avidan, S.; Brostow, G.; Ciss\u00e9, M.; Farinella, G. M.; Hassner, T. Eds. Springer Cham, 640\u2013656, 2022."},{"key":"379_CR7","series-title":"Lecture Notes in Computer Science","doi-asserted-by":"publisher","first-page":"341","DOI":"10.1007\/978-3-030-58542-6_21","volume-title":"Computer Vision\u2013ECCV 2020","author":"K Y Zhou","year":"2020","unstructured":"Zhou, K. Y.; Bhatnagar, B. L.; Pons-Moll, G. Unsupervised shape and pose disentanglement for 3D meshes. In: Computer Vision\u2013ECCV 2020. Lecture Notes in Computer Science, Vol. 12367. Vedaldi, A.; Bischof, H.; Brox, T.; Frahm, J. M. Eds. Springer Cham, 341\u2013357, 2020."},{"key":"379_CR8","series-title":"Lecture Notes in Computer Science","doi-asserted-by":"publisher","first-page":"19","DOI":"10.1007\/978-3-030-58580-8_2","volume-title":"Computer Vision\u2013ECCV 2020","author":"L Cosmo","year":"2020","unstructured":"Cosmo, L.; Norelli, A.; Halimi, O.; Kimmel, R.; Rodol\u00e0, E. LIMP: Learning latent shape representations with metric preservation priors. In: Computer Vision\u2013ECCV 2020. Lecture Notes in Computer Science, Vol. 12348. Vedaldi, A.; Bischof, H.; Brox, T.; Frahm, J. M. Eds. Springer Cham, 19\u201335, 2020."},{"key":"379_CR9","doi-asserted-by":"crossref","unstructured":"Huang, X.; Belongie, S. Arbitrary style transfer in real-time with adaptive instance normalization. In: Proceedings of the IEEE International Conference on Computer Vision, 1510\u20131519, 2017.","DOI":"10.1109\/ICCV.2017.167"},{"key":"379_CR10","doi-asserted-by":"crossref","unstructured":"Park, T.; Liu, M. Y.; Wang, T. C.; Zhu, J. Y. Semantic image synthesis with spatially-adaptive normalization. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 2332\u20132341, 2019.","DOI":"10.1109\/CVPR.2019.00244"},{"key":"379_CR11","unstructured":"Ulyanov, D.; Vedaldi, A.; Lempitsky, V. Instance normalization: The missing ingredient for fast stylization. arXiv preprint arXiv:1607.08022, 2016."},{"key":"379_CR12","series-title":"Lecture Notes in Computer Science","doi-asserted-by":"publisher","first-page":"176","DOI":"10.1007\/978-3-030-37731-1_15","volume-title":"MultiMedia Modeling","author":"Y G Chen","year":"2020","unstructured":"Chen, Y. G.; Chen, M. C.; Song, C. Y.; Ni, B. B. CartoonRenderer: An instance-based multi-style cartoon image translator. In: MultiMedia Modeling. Lecture Notes in Computer Science, Vol. 11961. Ro, Y., et al. Eds. Springer Cham, 176\u2013187, 2020."},{"key":"379_CR13","doi-asserted-by":"crossref","unstructured":"Wang, J. S.; Wen, C.; Fu, Y. W.; Lin, H. T.; Zou, T. Y.; Xue, X. Y.; Zhang, Y. D. Neural pose transfer by spatially adaptive instance normalization. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 5830\u20135838, 2020.","DOI":"10.1109\/CVPR42600.2020.00587"},{"issue":"1","key":"379_CR14","doi-asserted-by":"publisher","first-page":"258","DOI":"10.1609\/aaai.v36i1.19901","volume":"36","author":"H Y Chen","year":"2022","unstructured":"Chen, H. Y.; Tang, H.; Yu, Z. T.; Sebe, N.; Zhao, G. Y. Geometry-contrastive transformer for generalized 3D pose transfer. Proceedings of the AAAI Conference on Artificial Intelligence Vol. 36, No. 1, 258\u2013266, 2022.","journal-title":"Proceedings of the AAAI Conference on Artificial Intelligence"},{"key":"379_CR15","unstructured":"Song, C.; Wei, J.; Li, R.; Liu, F.; Lin, G. 3D pose transfer with correspondence learning and mesh refinement. In: Proceedings of the Advances in Neural Information Processing Systems, Vol. 34, 2021."},{"issue":"8","key":"379_CR16","doi-asserted-by":"publisher","first-page":"10488","DOI":"10.1109\/TPAMI.2023.3259059","volume":"45","author":"C Y Song","year":"2023","unstructured":"Song, C. Y.; Wei, J. C.; Li, R. B.; Liu, F. Y.; Lin, G. S. Unsupervised 3D pose transfer with cross consistency and dual reconstruction. IEEE Transactions on Pattern Analysis and Machine Intelligence Vol. 45, No. 8, 10488\u201310499, 2023.","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"issue":"3","key":"379_CR17","doi-asserted-by":"publisher","first-page":"331","DOI":"10.1007\/s41095-022-0271-y","volume":"8","author":"M H Guo","year":"2022","unstructured":"Guo, M. H.; Xu, T. X.; Liu, J. J.; Liu, Z. N.; Jiang, P. T.; Mu, T. J.; Zhang, S. H.; Martin, R. R.; Cheng, M. M.; Hu, S. M. Attention mechanisms in computer vision: A survey. Computational Visual Media Vol. 8, No. 3, 331\u2013368, 2022.","journal-title":"Computational Visual Media"},{"issue":"1","key":"379_CR18","doi-asserted-by":"publisher","first-page":"33","DOI":"10.1007\/s41095-021-0247-3","volume":"8","author":"Y F Xu","year":"2022","unstructured":"Xu, Y. F.; Wei, H. P.; Lin, M. X.; Deng, Y. Y.; Sheng, K. K.; Zhang, M. D.; Tang, F.; Dong, W. M.; Huang, F. Y.; Xu, C. S. Transformers in computational visual media: A survey. Computational Visual Media Vol. 8, No. 1, 33\u201362, 2022.","journal-title":"Computational Visual Media"},{"key":"379_CR19","doi-asserted-by":"crossref","unstructured":"Sumner, R. W.; Popovi\u0107 J. Deformation transfer for triangle meshes. In: Proceedings of the ACM SIGGRAPH Papers, 399\u2013405, 2004.","DOI":"10.1145\/1186562.1015736"},{"issue":"3","key":"379_CR20","doi-asserted-by":"publisher","first-page":"84","DOI":"10.1145\/1276377.1276482","volume":"26","author":"W W Xu","year":"2007","unstructured":"Xu, W. W.; Zhou, K.; Yu, Y. Z.; Tan, Q. F.; Peng, Q. S.; Guo, B. N. Gradient domain editing of deforming mesh sequences. ACM Transactions on Graphics Vol. 26, No. 3, 84\u2013es, 2007.","journal-title":"ACM Transactions on Graphics"},{"key":"379_CR21","doi-asserted-by":"crossref","unstructured":"Domadiya, P. M.; Shah, D. P.; Mitra, S. Guided deformation transfer. In: Proceedings of the 16th ACM SIGGRAPH European Conference on Visual Media Production, Article No. 7, 2019.","DOI":"10.1145\/3359998.3369408"},{"key":"379_CR22","doi-asserted-by":"publisher","first-page":"11","DOI":"10.1016\/j.cag.2020.04.002","volume":"89","author":"J Basset","year":"2020","unstructured":"Basset, J.; Wuhrer, S.; Boyer, E.; Multon, F. Contact preserving shape transfer: Retargeting motion from one shape to another. Computers & Graphics Vol. 89, 11\u201323, 2020.","journal-title":"Computers & Graphics"},{"key":"379_CR23","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1016\/j.gmod.2018.05.003","volume":"98","author":"J Yang","year":"2018","unstructured":"Yang, J.; Gao, L.; Lai, Y. K.; Rosin, P. L.; Xia, S. H. Biharmonic deformation transfer with automatic key point selection. Graphical Models Vol. 98, 1\u201313, 2018.","journal-title":"Graphical Models"},{"key":"379_CR24","doi-asserted-by":"crossref","unstructured":"Ben-Chen, M.; Weber, O.; Gotsman, C. Variational harmonic maps for space deformation. ACM Transactions on Graphics Vol. 28, No. 3, Article No. 34, 2009.","DOI":"10.1145\/1531326.1531340"},{"key":"379_CR25","doi-asserted-by":"crossref","unstructured":"Jacobson, A.; Baran, I.; Popovi\u0107 J.; Sorkine, O. Bounded biharmonic weights for real-time deformation. ACM Transactions on Graphics Vol. 30, No. 4, Article No. 78, 2011.","DOI":"10.1145\/2010324.1964973"},{"key":"379_CR26","doi-asserted-by":"crossref","unstructured":"Baran, I.; Vlasic, D.; Grinspun, E.; Popovi\u0107 J. Semantic deformation transfer. ACM Transactions on Graphics Vol. 28, No. 3, Article No. 36, 2009.","DOI":"10.1145\/1531326.1531342"},{"key":"379_CR27","unstructured":"Chen, H.; Tang, H.; Sebe, N.; Zhao, G. AniFormer: Datadriven 3D animation with transformer. In: Proceedings of the British Machine Vision Conference, 2021."},{"key":"379_CR28","doi-asserted-by":"crossref","unstructured":"Gao, L.; Yang, J.; Qiao, Y. L.; Lai, Y. K.; Rosin, P. L.; Xu, W. W.; Xia, S. H. Automatic unpaired shape deformation transfer. ACM Transactions on Graphics Vol. 37, No. 6, Article No. 237, 2018.","DOI":"10.1145\/3272127.3275028"},{"key":"379_CR29","doi-asserted-by":"crossref","unstructured":"Chen, H. Y.; Tang, H.; Shi, H. L.; Peng, W.; Sebe, N.; Zhao, G. Y. Intrinsic-extrinsic preserved GANs for unsupervised 3D pose transfer. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, 8610\u20138619, 2021.","DOI":"10.1109\/ICCV48922.2021.00851"},{"key":"379_CR30","doi-asserted-by":"crossref","unstructured":"Wang, Y. F.; Aigerman, N.; Kim, V. G.; Chaudhuri, S.; Sorkine-Hornung, O. Neural cages for detail-preserving 3D deformations. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 72\u201380, 2020.","DOI":"10.1109\/CVPR42600.2020.00015"},{"key":"379_CR31","unstructured":"Vaswani, A.; Shazeer, N.; Parmar, N.; Uszkoreit, J.; Jones, L.; Gomez, A. N.; Kaiser, L.; Polosukhin, I. Attention is all you need. In: Proceedings of the 31st International Conference on Neural Information Processing Systems, 6000\u20136010, 2017."},{"key":"379_CR32","unstructured":"Dosovitskiy, A.; Beyer, L.; Kolesnikov, A.; Weissenborn, D.; Zhai, X. H.; Unterthiner, T.; Dehghani, M.; Minderer, M.; Heigold, G.; Gelly, S.; et al. An image is worth 16x16 words: Transformers for image recognition at scale. arXiv preprint arXiv:2010.11929, 2020."},{"key":"379_CR33","doi-asserted-by":"crossref","unstructured":"Lin, K.; Wang, L. J.; Liu, Z. C. End-to-end human pose and mesh reconstruction with transformers. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, 1954\u20131963, 2021.","DOI":"10.1109\/CVPR46437.2021.00199"},{"key":"379_CR34","doi-asserted-by":"crossref","unstructured":"Lin, K.; Wang, L. J.; Liu, Z. C. Mesh graphormer. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, 12919\u201312928, 2021.","DOI":"10.1109\/ICCV48922.2021.01270"},{"key":"379_CR35","doi-asserted-by":"crossref","unstructured":"Misra, I.; Girdhar, R.; Joulin, A. An end-to-end transformer model for 3D object detection. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, 2886\u20132897, 2021.","DOI":"10.1109\/ICCV48922.2021.00290"},{"key":"379_CR36","doi-asserted-by":"crossref","unstructured":"Mao, J. G.; Xue, Y. J.; Niu, M. Z.; Bai, H. Y.; Feng, J. S.; Liang, X. D.; Xu, H.; Xu, C. J. Voxel transformer for 3D object detection. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, 3144\u20133153, 2021.","DOI":"10.1109\/ICCV48922.2021.00315"},{"key":"379_CR37","first-page":"20014","volume":"34","author":"A Ali","year":"2021","unstructured":"Ali, A.; Touvron, H.; Caron, M.; Bojanowski, P.; Douze, M.; Joulin, A.; Laptev, I.; Neverova, N.; Synnaeve, G.; Verbeek, J.; et al. Xcit: Cross-covariance image transformers. In: Proceedings of the Advances in Neural Information Processing Systems, Vol. 34, 20014\u201320027, 2021.","journal-title":"Proceedings of the Advances in Neural Information Processing Systems"},{"issue":"2","key":"379_CR38","doi-asserted-by":"publisher","first-page":"195","DOI":"10.1111\/cgf.14468","volume":"41","author":"P Chandran","year":"2022","unstructured":"Chandran, P.; Zoss, G.; Gross, M.; Gotardo, P.; Bradley, D. Shape transformers: Topology-independent 3D shape models using transformers. Computer Graphics Forum Vol. 41, No. 2, 195\u2013207, 2022.","journal-title":"Computer Graphics Forum"},{"key":"379_CR39","first-page":"851","volume":"2","author":"M Loper","year":"2023","unstructured":"Loper, M.; Mahmood, N.; Romero, J.; Pons-Moll, G.; Black, M. J. SMPL: A skinned multi-person linear model. Seminal Graphics Papers: Pushing the Boundaries Vol. 2, Article No. 88, 851\u2013866, 2023.","journal-title":"Seminal Graphics Papers: Pushing the Boundaries"},{"key":"379_CR40","doi-asserted-by":"crossref","unstructured":"Bogo, F.; Romero, J.; Loper, M.; Black, M. J. FAUST: Dataset and evaluation for 3D mesh registration. In:Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, 3794\u20133801, 2014.","DOI":"10.1109\/CVPR.2014.491"},{"key":"379_CR41","doi-asserted-by":"crossref","unstructured":"Bhatnagar, B.; Tiwari, G.; Theobalt, C.; Pons-Moll, G. Multi-garment net: Learning to dress 3D people from images. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, 5419\u20135429, 2019.","DOI":"10.1109\/ICCV.2019.00552"},{"key":"379_CR42","doi-asserted-by":"crossref","unstructured":"Zuffi, S.; Kanazawa, A.; Jacobs, D. W.; Black, M. J. 3D menagerie: Modeling the 3D shape and pose of animals. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, 5524\u20135532, 2017.","DOI":"10.1109\/CVPR.2017.586"},{"key":"379_CR43","unstructured":"Loshchilov, I.; Hutter, F. Decoupled weight decay regularization. arXiv preprint arXiv:1711.05101, 2017."},{"key":"379_CR44","unstructured":"Paszke, A.; Gross, S.; Massa, F.; Lerer, A.; Bradbury, J.; Chanan, G.; Killeen, T.; Lin, Z.; Gimelshein, N.; Antiga, L.; et al. PyTorch: An imperative style, high-performance deep learning library. In: Proceedings of the 33rd International Conference on Neural Information Processing Systems, Article No. 721, 8026\u20138037, 2019."},{"key":"379_CR45","doi-asserted-by":"crossref","unstructured":"Fan, H. Q.; Su, H.; Guibas, L. A point set generation network for 3D object reconstruction from a single image. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, 2463\u20132471, 2017.","DOI":"10.1109\/CVPR.2017.264"},{"key":"379_CR46","doi-asserted-by":"crossref","unstructured":"Mahmood, N.; Ghorbani, N.; Troje, N. F.; Pons-Moll, G.; Black, M. AMASS: Archive of motion capture as surface shapes. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, 5441\u20135450, 2019.","DOI":"10.1109\/ICCV.2019.00554"}],"container-title":["Computational Visual Media"],"original-title":[],"link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s41095-023-0379-8.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s41095-023-0379-8\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s41095-023-0379-8.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"},{"URL":"http:\/\/xplorestaging.ieee.org\/ielx8\/10750449\/10884987\/10884991.pdf?arnumber=10884991","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,11,5]],"date-time":"2025-11-05T18:38:15Z","timestamp":1762367895000},"score":1,"resource":{"primary":{"URL":"https:\/\/ieeexplore.ieee.org\/document\/10884991\/"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2024,12]]},"references-count":46,"journal-issue":{"issue":"6"},"URL":"https:\/\/doi.org\/10.1007\/s41095-023-0379-8","relation":{},"ISSN":["2096-0662","2096-0433"],"issn-type":[{"type":"electronic","value":"2096-0662"},{"type":"print","value":"2096-0433"}],"subject":[],"published":{"date-parts":[[2024,12]]},"assertion":[{"value":"13 March 2023","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"9 September 2023","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"22 March 2024","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"The authors have no competing interests to declare that are relevant to the content of this article.","order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declaration of competing interest"}}]}}