{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,8,17]],"date-time":"2026-08-17T15:01:46Z","timestamp":1786978906183,"version":"build-2736575974"},"publisher-location":"Cham","reference-count":56,"publisher":"Springer Nature Switzerland","isbn-type":[{"value":"9783031729126","type":"print"},{"value":"9783031729133","type":"electronic"}],"license":[{"start":{"date-parts":[[2024,12,2]],"date-time":"2024-12-02T00:00:00Z","timestamp":1733097600000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2024,12,2]],"date-time":"2024-12-02T00:00:00Z","timestamp":1733097600000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":[],"published-print":{"date-parts":[[2025]]},"DOI":"10.1007\/978-3-031-72913-3_8","type":"book-chapter","created":{"date-parts":[[2024,12,1]],"date-time":"2024-12-01T16:46:35Z","timestamp":1733071595000},"page":"129-145","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":5,"title":["SHIC: Shape-Image Correspondences with\u00a0No Keypoint Supervision"],"prefix":"10.1007","author":[{"given":"Aleksandar","family":"Shtedritski","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Christian","family":"Rupprecht","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Andrea","family":"Vedaldi","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2024,12,2]]},"reference":[{"key":"8_CR1","unstructured":"Amir, S., Gandelsman, Y., Bagon, S., Dekel, T.: Deep ViT features as dense visual descriptors. CoRR abs\/2112.05814 (2021)"},{"key":"8_CR2","doi-asserted-by":"crossref","unstructured":"Bourdev, L.D., Malik, J.: Poselets: body part detectors trained using 3D human pose annotations. In: Proceedings of ICCV (2009)","DOI":"10.1109\/ICCV.2009.5459303"},{"key":"8_CR3","doi-asserted-by":"crossref","unstructured":"Cao, Z., Simon, T., Wei, S., Sheikh, Y.: Realtime multi-person 2D pose estimation using part affinity fields. In: Proceedings of CVPR (2017)","DOI":"10.1109\/CVPR.2017.143"},{"key":"8_CR4","doi-asserted-by":"crossref","unstructured":"Caron, M., et al.: Emerging properties in self-supervised vision transformers. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 9650\u20139660 (2021)","DOI":"10.1109\/ICCV48922.2021.00951"},{"key":"8_CR5","unstructured":"Chen, J., Wang, L., Li, X., Fang, Y.: Arbicon-net: arbitrary continuous geometric transformation networks for image registration. In: Wallach, H., Larochelle, H., Beygelzimer, A., d\u2019 Alch\u00e9-Buc, F., Fox, E., Garnett, R. (eds.) Advances in Neural Information Processing Systems, vol.\u00a032. Curran Associates, Inc. (2019)"},{"key":"8_CR6","doi-asserted-by":"crossref","unstructured":"Chen, R., Chen, Y., Jiao, N., Jia, K.: Fantasia3d: disentangling geometry and appearance for high-quality text-to-3D content creation. arXiv.cs abs\/2303.13873 (2023)","DOI":"10.1109\/ICCV51070.2023.02033"},{"key":"8_CR7","doi-asserted-by":"crossref","unstructured":"Dutt, N.S., Muralikrishnan, S., Mitra, N.J.: Diffusion 3D features (Diff3F): decorating untextured shapes with distilled semantic features. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR), pp. 4494\u20134504 (2024)","DOI":"10.1109\/CVPR52733.2024.00430"},{"key":"8_CR8","doi-asserted-by":"crossref","unstructured":"Felzenszwalb, P.F., McAllester, D.A., Ramanan, D.: A discriminatively trained, multiscale, deformable part model. In: Proceedings of CVPR (2008)","DOI":"10.1109\/CVPR.2008.4587597"},{"key":"8_CR9","doi-asserted-by":"crossref","unstructured":"G\u00fcler, R.A., Neverova, N., Kokkinos, I.: Densepose: dense human pose estimation in the wild. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 7297\u20137306 (2018)","DOI":"10.1109\/CVPR.2018.00762"},{"key":"8_CR10","doi-asserted-by":"crossref","unstructured":"G\u00fcler, R.A., Neverova, N., Kokkinos, I.: DensePose: dense human pose estimation in the wild. In: Proceedings of CVPR (2018)","DOI":"10.1109\/CVPR.2018.00762"},{"key":"8_CR11","doi-asserted-by":"crossref","unstructured":"Ham, B., Cho, M., Schmid, C., Ponce, J.: Proposal flow. In: Proceedings of CVPR (2016)","DOI":"10.1109\/CVPR.2016.378"},{"issue":"7","key":"8_CR12","doi-asserted-by":"publisher","first-page":"1711","DOI":"10.1109\/TPAMI.2017.2724510","volume":"40","author":"B Ham","year":"2017","unstructured":"Ham, B., Cho, M., Schmid, C., Ponce, J.: Proposal flow: semantic correspondences from object proposals. IEEE Trans. Pattern Anal. Mach. Intell. 40(7), 1711\u20131725 (2017)","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"8_CR13","unstructured":"Hedlin, E., et al.: Unsupervised semantic correspondence using stable diffusion. arXiv.cs (2023)"},{"key":"8_CR14","doi-asserted-by":"crossref","unstructured":"Jeon, S., Kim, S., Min, D., Sohn, K.: PARN: pyramidal affine regression networks for dense semantic correspondence. In: Proceedings of the European Conference on Computer Vision (ECCV), pp. 351\u2013366 (2018)","DOI":"10.1007\/978-3-030-01231-1_22"},{"key":"8_CR15","doi-asserted-by":"crossref","unstructured":"Kanazawa, A., Jacobs, D.W., Chandraker, M.: WarpNet: weakly supervised matching for single-view reconstruction. In: Proceedings of CVPR (2016)","DOI":"10.1109\/CVPR.2016.354"},{"key":"8_CR16","doi-asserted-by":"crossref","unstructured":"Kreiss, S., Bertoni, L., Alahi, A.: Pifpaf: composite fields for human pose estimation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 11977\u201311986 (2019)","DOI":"10.1109\/CVPR.2019.01225"},{"key":"8_CR17","doi-asserted-by":"crossref","unstructured":"Kulkarni, N., Gupta, A., Fouhey, D.F., Tulsiani, S.: Articulation-aware canonical surface mapping. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 452\u2013461 (2020)","DOI":"10.1109\/CVPR42600.2020.00053"},{"key":"8_CR18","doi-asserted-by":"crossref","unstructured":"Kulkarni, N., Gupta, A., Tulsiani, S.: Canonical surface mapping via geometric cycle consistency. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 2202\u20132211 (2019)","DOI":"10.1109\/ICCV.2019.00229"},{"key":"8_CR19","doi-asserted-by":"crossref","unstructured":"Li, X., Lu, J., Han, K., Prisacariu, V.: SD4Match: learning to prompt stable diffusion model for semantic matching. arXiv preprint arXiv:2310.17569 (2023)","DOI":"10.1109\/CVPR52733.2024.02602"},{"key":"8_CR20","doi-asserted-by":"crossref","unstructured":"Liu, S., et\u00a0al.: Grounding dino: marrying dino with grounded pre-training for open-set object detection. arXiv preprint arXiv:2303.05499 (2023)","DOI":"10.1007\/978-3-031-72970-6_3"},{"key":"8_CR21","unstructured":"Luo, G., Dunlap, L., Park, D.H., Holynski, A., Darrell, T.: Diffusion hyperfeatures: searching through time and space for semantic correspondence. In: Advances in Neural Information Processing Systems (2023)"},{"key":"8_CR22","doi-asserted-by":"crossref","unstructured":"Melekhov, I., Tiulpin, A., Sattler, T., Pollefeys, M., Rahtu, E., Kannala, J.: DGC-Net: dense geometric correspondence network. In: 2019 IEEE Winter Conference on Applications of Computer Vision (WACV), pp. 1034\u20131042. IEEE (2019)","DOI":"10.1109\/WACV.2019.00115"},{"key":"8_CR23","doi-asserted-by":"crossref","unstructured":"Morreale, L., Aigerman, N., Kim, V.G., Mitra, N.J.: Neural semantic surface maps. In: Computer Graphics Forum, vol.\u00a043, p. e15005. Wiley Online Library (2024)","DOI":"10.1111\/cgf.15005"},{"key":"8_CR24","first-page":"17258","volume":"33","author":"N Neverova","year":"2020","unstructured":"Neverova, N., Novotny, D., Szafraniec, M., Khalidov, V., Labatut, P., Vedaldi, A.: Continuous surface embeddings. Adv. Neural. Inf. Process. Syst. 33, 17258\u201317270 (2020)","journal-title":"Adv. Neural. Inf. Process. Syst."},{"key":"8_CR25","doi-asserted-by":"crossref","unstructured":"Neverova, N., Sanakoyeu, A., Labatut, P., Novotny, D., Vedaldi, A.: Discovering relationships between object categories via universal canonical maps. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 404\u2013413 (2021)","DOI":"10.1109\/CVPR46437.2021.00047"},{"key":"8_CR26","doi-asserted-by":"crossref","unstructured":"Newell, A., Yang, K., Deng, J.: Stacked hourglass networks for human pose estimation. In: Proceedings of ECCV (2016)","DOI":"10.1007\/978-3-319-46484-8_29"},{"key":"8_CR27","unstructured":"OpenAI: Chatgpt. https:\/\/chat.openai.com\/"},{"key":"8_CR28","unstructured":"Oquab, M., et\u00a0al.: Dinov2: learning robust visual features without supervision. arXiv preprint arXiv:2304.07193 (2023)"},{"key":"8_CR29","doi-asserted-by":"crossref","unstructured":"Peebles, W., Zhu, J.Y., Zhang, R., Torralba, A., Efros, A.A., Shechtman, E.: GAN-supervised dense visual alignment. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 13470\u201313481 (2022)","DOI":"10.1109\/CVPR52688.2022.01311"},{"key":"8_CR30","doi-asserted-by":"crossref","unstructured":"Pereira, T., et al.: Fast animal pose estimation using deep neural networks. bioRxiv (2018)","DOI":"10.1101\/331181"},{"key":"8_CR31","unstructured":"Radford, A., et\u00a0al.: Learning transferable visual models from natural language supervision. In: International Conference on Machine Learning, pp. 8748\u20138763. PMLR (2021)"},{"key":"8_CR32","unstructured":"Ramesh, A., Dhariwal, P., Nichol, A., Chu, C., Chen, M.: Hierarchical text-conditional image generation with clip latents. arXiv preprint arXiv:2204.06125 (2022)"},{"key":"8_CR33","doi-asserted-by":"crossref","unstructured":"Rempe, D., Birdal, T., Hertzmann, A., Yang, J., Sridhar, S., Guibas, L.J.: Humor: 3D human motion model for robust pose estimation. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 11488\u201311499 (2021)","DOI":"10.1109\/ICCV48922.2021.01129"},{"key":"8_CR34","doi-asserted-by":"crossref","unstructured":"Rocco, I., Arandjelovic, R., Sivic, J.: Convolutional neural network architecture for geometric matching. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 6148\u20136157 (2017)","DOI":"10.1109\/CVPR.2017.12"},{"key":"8_CR35","doi-asserted-by":"crossref","unstructured":"Rombach, R., Blattmann, A., Lorenz, D., Esser, P., Ommer, B.: High-resolution image synthesis with latent diffusion models (2021)","DOI":"10.1109\/CVPR52688.2022.01042"},{"key":"8_CR36","doi-asserted-by":"crossref","unstructured":"Shtedritski, A., Rupprecht, C., Vedaldi, A.: Learning universal semantic correspondences with no supervision and automatic data curation. In: Proceedings of IEEE\/CVF International Conference on Computer Vision (ICCV) Workshops (2023)","DOI":"10.1109\/ICCVW60793.2023.00100"},{"key":"8_CR37","unstructured":"Tang, L., Jia, M., Wang, Q., Phoo, C.P., Hariharan, B.: Emergent correspondence from image diffusion. In: Thirty-Seventh Conference on Neural Information Processing Systems (2023)"},{"key":"8_CR38","unstructured":"Thewlis, J., Bilen, H., Vedaldi, A.: Unsupervised learning of object frames by dense equivariant image labelling. In: Proceedings of Advances in Neural Information Processing Systems (NeurIPS) (2017)"},{"key":"8_CR39","unstructured":"Thewlis, J., Bilen, H., Vedaldi, A.: Modelling and unsupervised learning of symmetric deformable object categories. In: Proceedings of Advances in Neural Information Processing Systems (NeurIPS) (2018)"},{"key":"8_CR40","first-page":"14278","volume":"33","author":"P Truong","year":"2020","unstructured":"Truong, P., Danelljan, M., Gool, L.V., Timofte, R.: GOCor: bringing globally optimized correspondence volumes into your neural network. Adv. Neural. Inf. Process. Syst. 33, 14278\u201314290 (2020)","journal-title":"Adv. Neural. Inf. Process. Syst."},{"key":"8_CR41","doi-asserted-by":"crossref","unstructured":"Truong, P., Danelljan, M., Timofte, R.: GLU-Net: global-local universal network for dense flow and correspondences. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 6258\u20136268 (2020)","DOI":"10.1109\/CVPR42600.2020.00629"},{"key":"8_CR42","doi-asserted-by":"crossref","unstructured":"Truong, P., Danelljan, M., Yu, F., Van\u00a0Gool, L.: Warp consistency for unsupervised learning of dense correspondences. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 10346\u201310356 (2021)","DOI":"10.1109\/ICCV48922.2021.01018"},{"key":"8_CR43","doi-asserted-by":"crossref","unstructured":"Truong, P., Danelljan, M., Yu, F., Van\u00a0Gool, L.: Probabilistic warp consistency for weakly-supervised semantic correspondences. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 8708\u20138718 (2022)","DOI":"10.1109\/CVPR52688.2022.00851"},{"key":"8_CR44","doi-asserted-by":"crossref","unstructured":"Waldmann, U., et al.: 3D-muppet: 3D multi-pigeon pose estimation and tracking. arXiv preprint arXiv:2308.15316 (2023)","DOI":"10.1007\/s11263-024-02074-y"},{"key":"8_CR45","doi-asserted-by":"crossref","unstructured":"Wei, S., Ramakrishna, V., Kanade, T., Sheikh, Y.: Convolutional pose machines. In: Proceedings of CVPR (2016)","DOI":"10.1109\/CVPR.2016.511"},{"key":"8_CR46","unstructured":"Wu, S., Jakab, T., Rupprecht, C., Vedaldi, A.: DOVE: learning deformable 3D objects by watching videos. arXiv (2021)"},{"key":"8_CR47","doi-asserted-by":"crossref","unstructured":"Wu, S., Li, R., Jakab, T., Rupprecht, C., Vedaldi, A.: MagicPony: learning articulated 3D animals in the wild. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (CVPR) (2023)","DOI":"10.1109\/CVPR52729.2023.00849"},{"key":"8_CR48","doi-asserted-by":"crossref","unstructured":"Yang, L., Kang, B., Huang, Z., Xu, X., Feng, J., Zhao, H.: Depth anything: unleashing the power of large-scale unlabeled data. arXiv:2401.10891 (2024)","DOI":"10.1109\/CVPR52733.2024.00987"},{"key":"8_CR49","doi-asserted-by":"crossref","unstructured":"Zhang, H., et al.: Pymaf: 3D human pose and shape regression with pyramidal mesh alignment feedback loop. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 11446\u201311456 (2021)","DOI":"10.1109\/ICCV48922.2021.01125"},{"key":"8_CR50","unstructured":"Zhang, J., et al.: A Tale of Two Features: Stable Diffusion Complements DINO for Zero-Shot Semantic Correspondence. arXiv preprint arxiv:2305.15347 (2023)"},{"key":"8_CR51","doi-asserted-by":"crossref","unstructured":"Zhang, J., et al.: Telling left from right: identifying geometry-aware semantic correspondence. arXiv.cs (2023)","DOI":"10.1109\/CVPR52733.2024.00297"},{"key":"8_CR52","doi-asserted-by":"crossref","unstructured":"Zhang, L., Rao, A., Agrawala, M.: Adding conditional control to text-to-image diffusion models (2023)","DOI":"10.1109\/ICCV51070.2023.00355"},{"key":"8_CR53","doi-asserted-by":"crossref","unstructured":"Zhang, N., Donahue, J., Girshick, R.B., Darrell, T.: Part-based R-CNNs for fine-grained category detection. In: Proceedings of ECCV (2014)","DOI":"10.1007\/978-3-319-10590-1_54"},{"key":"8_CR54","doi-asserted-by":"crossref","unstructured":"Zuffi, S., Kanazawa, A., Berger-Wolf, T., Black, M.J.: Three-d safari: learning to estimate zebra pose, shape, and texture from images \u201cin the wild\u201d. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 5359\u20135368 (2019)","DOI":"10.1109\/ICCV.2019.00546"},{"key":"8_CR55","doi-asserted-by":"crossref","unstructured":"Zuffi, S., Kanazawa, A., Black, M.J.: Lions and tigers and bears: capturing non-rigid, 3D, articulated shape from images. In: 2018 IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 3955\u20133963 (2018)","DOI":"10.1109\/CVPR.2018.00416"},{"key":"8_CR56","doi-asserted-by":"crossref","unstructured":"Zuffi, S., Kanazawa, A., Jacobs, D.W., Black, M.J.: 3D menagerie: modeling the 3D shape and pose of animals. In: Proceedings of CVPR (2017)","DOI":"10.1109\/CVPR.2017.586"}],"container-title":["Lecture Notes in Computer Science","Computer Vision \u2013 ECCV 2024"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/978-3-031-72913-3_8","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2024,12,1]],"date-time":"2024-12-01T18:22:54Z","timestamp":1733077374000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/978-3-031-72913-3_8"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2024,12,2]]},"ISBN":["9783031729126","9783031729133"],"references-count":56,"URL":"https:\/\/doi.org\/10.1007\/978-3-031-72913-3_8","relation":{},"ISSN":["0302-9743","1611-3349"],"issn-type":[{"value":"0302-9743","type":"print"},{"value":"1611-3349","type":"electronic"}],"subject":[],"published":{"date-parts":[[2024,12,2]]},"assertion":[{"value":"2 December 2024","order":1,"name":"first_online","label":"First Online","group":{"name":"ChapterHistory","label":"Chapter History"}},{"value":"ECCV","order":1,"name":"conference_acronym","label":"Conference Acronym","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"European Conference on Computer Vision","order":2,"name":"conference_name","label":"Conference Name","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"Milan","order":3,"name":"conference_city","label":"Conference City","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"Italy","order":4,"name":"conference_country","label":"Conference Country","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"2024","order":5,"name":"conference_year","label":"Conference Year","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"29 September 2024","order":7,"name":"conference_start_date","label":"Conference Start Date","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"4 October 2024","order":8,"name":"conference_end_date","label":"Conference End Date","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"18","order":9,"name":"conference_number","label":"Conference Number","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"eccv2024","order":10,"name":"conference_id","label":"Conference ID","group":{"name":"ConferenceInfo","label":"Conference Information"}},{"value":"https:\/\/eccv2024.ecva.net\/","order":11,"name":"conference_url","label":"Conference URL","group":{"name":"ConferenceInfo","label":"Conference Information"}}]}}