{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,3,11]],"date-time":"2026-03-11T01:53:50Z","timestamp":1773194030188,"version":"3.50.1"},"reference-count":34,"publisher":"Springer Science and Business Media LLC","issue":"32","license":[{"start":{"date-parts":[[2024,2,27]],"date-time":"2024-02-27T00:00:00Z","timestamp":1708992000000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2024,2,27]],"date-time":"2024-02-27T00:00:00Z","timestamp":1708992000000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100003725","name":"National Research Foundation of Korea","doi-asserted-by":"crossref","award":["BK21 FOUR"],"award-info":[{"award-number":["BK21 FOUR"]}],"id":[{"id":"10.13039\/501100003725","id-type":"DOI","asserted-by":"crossref"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Multimed Tools Appl"],"DOI":"10.1007\/s11042-024-18653-7","type":"journal-article","created":{"date-parts":[[2024,2,27]],"date-time":"2024-02-27T05:01:57Z","timestamp":1709010117000},"page":"78577-78592","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":2,"title":["Multimodal contrastive learning using point clouds and their rendered images"],"prefix":"10.1007","volume":"83","author":[{"given":"Wonyong","family":"Lee","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0001-9013-2338","authenticated-orcid":false,"given":"Hyungki","family":"Kim","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2024,2,27]]},"reference":[{"key":"18653_CR1","doi-asserted-by":"crossref","unstructured":"Lin C-H, Kong C, Lucey S (2018) Learning efficient point cloud generation for dense 3D object reconstruction. In: Proceedings of the Thirty-Second AAAI Conference on Artificial Intelligence and Thirtieth Innovative Applications of Artificial Intelligence Conference and Eighth AAAI Symposium on Educational Advances in Artificial Intelligence. AAAI Press, New Orleans, pp 7114\u20137121","DOI":"10.1609\/aaai.v32i1.12278"},{"key":"18653_CR2","doi-asserted-by":"publisher","first-page":"187","DOI":"10.1007\/s41095-021-0229-5","volume":"7","author":"M-H Guo","year":"2021","unstructured":"Guo M-H, Cai J-X, Liu Z-N et al (2021) PCT: point cloud transformer. Comp Visual Media 7:187\u2013199. https:\/\/doi.org\/10.1007\/s41095-021-0229-5","journal-title":"Comp Visual Media"},{"key":"18653_CR3","doi-asserted-by":"publisher","unstructured":"Charles RQ, Su H, Kaichun M, Guibas LJ (2017) PointNet: deep learning on point sets for 3D classification and segmentation. In: 2017 IEEE Conference on Computer Vision and Pattern Recognition (CVPR). IEEE, Honolulu, HI, pp 77\u201385. https:\/\/doi.org\/10.1109\/CVPR.2017.16","DOI":"10.1109\/CVPR.2017.16"},{"issue":"146","key":"18653_CR4","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1145\/3326362","volume":"38","author":"Y Wang","year":"2019","unstructured":"Wang Y, Sun Y, Liu Z et al (2019) Dynamic graph CNN for learning on point clouds. ACM Trans Graph 38(146):1\u2013146. https:\/\/doi.org\/10.1145\/3326362","journal-title":"ACM Trans Graph"},{"key":"18653_CR5","doi-asserted-by":"publisher","unstructured":"Zhang Z, Girdhar R, Joulin A, Misra I (2021) Self-Supervised Pretraining of 3D Features on any Point-Cloud. In: 2021 IEEE\/CVF International Conference on Computer Vision (ICCV). IEEE, Montreal, QC, Canada, pp 10232\u201310243. https:\/\/doi.org\/10.1109\/ICCV48922.2021.01009","DOI":"10.1109\/ICCV48922.2021.01009"},{"key":"18653_CR6","doi-asserted-by":"publisher","unstructured":"Afham M, Dissanayake I, Dissanayake D, et al (2022) CrossPoint: self-supervised cross-modal contrastive learning for 3D point cloud understanding. In: 2022 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR). IEEE, New Orleans, LA, USA, pp 9892\u20139902. https:\/\/doi.org\/10.1109\/CVPR52688.2022.00967","DOI":"10.1109\/CVPR52688.2022.00967"},{"key":"18653_CR7","doi-asserted-by":"publisher","unstructured":"Huang S, Xie Y, Zhu S-C, Zhu Y (2021) Spatio-temporal Self-Supervised Representation Learning for 3D Point Clouds. In: 2021 IEEE\/CVF International Conference on Computer Vision (ICCV). IEEE, Montreal, QC, Canada, pp 6515\u20136525. https:\/\/doi.org\/10.1109\/ICCV48922.2021.00647","DOI":"10.1109\/ICCV48922.2021.00647"},{"key":"18653_CR8","doi-asserted-by":"publisher","unstructured":"Du B, Gao X, Hu W, Li X (2021) Self-contrastive learning with hard negative sampling for self-supervised point cloud learning. In: Proceedings of the 29th ACM International Conference on Multimedia. Association for Computing Machinery, New York, NY, USA, pp 3133\u20133142. https:\/\/doi.org\/10.1145\/3474085.3475458","DOI":"10.1145\/3474085.3475458"},{"key":"18653_CR9","doi-asserted-by":"publisher","unstructured":"Xie S, Gu J, Guo D et al (2020) PointContrast: unsupervised pre-training for 3D point cloud understanding. In: Vedaldi A, Bischof H, Brox T, Frahm J-M (eds) Computer vision \u2013 ECCV 2020. Springer International Publishing, Cham, pp 574\u2013591. https:\/\/doi.org\/10.1007\/978-3-030-58580-8_34","DOI":"10.1007\/978-3-030-58580-8_34"},{"key":"18653_CR10","doi-asserted-by":"publisher","unstructured":"Chang AX, Funkhouser T, Guibas L, Hanrahan P, Huang Q, Li Z, Savarese S, Yu F (2015) Shapenet: an information-rich 3d model repository. arXiv preprint arXiv:1512.03012. https:\/\/doi.org\/10.48550\/arXiv.1512.03012","DOI":"10.48550\/arXiv.1512.03012"},{"key":"18653_CR11","doi-asserted-by":"publisher","unstructured":"Su H, Maji S, Kalogerakis E, Learned-Miller E (2015) Multi-view Convolutional Neural Networks for 3D Shape Recognition. In: 2015 IEEE International Conference on Computer Vision (ICCV). IEEE, Santiago, Chile, pp 945\u2013953. https:\/\/doi.org\/10.1109\/ICCV.2015.114","DOI":"10.1109\/ICCV.2015.114"},{"key":"18653_CR12","doi-asserted-by":"publisher","unstructured":"Pang G, Neumann U (2016) 3D point cloud object detection with multi-view convolutional neural network. In: 2016 23rd International Conference on Pattern Recognition (ICPR). pp 585\u2013590. https:\/\/doi.org\/10.1109\/ICPR.2016.7899697","DOI":"10.1109\/ICPR.2016.7899697"},{"key":"18653_CR13","doi-asserted-by":"publisher","unstructured":"He K, Zhang X, Ren S, Sun J (2016) Deep residual learning for image recognition. In: 2016 IEEE Conference on Computer Vision and Pattern Recognition (CVPR). IEEE, Las Vegas, NV, USA, pp 770\u2013778. https:\/\/doi.org\/10.1109\/CVPR.2016.90","DOI":"10.1109\/CVPR.2016.90"},{"key":"18653_CR14","doi-asserted-by":"publisher","unstructured":"Maturana D, Scherer S (2015) VoxNet: A 3D convolutional neural network for real-time object recognition. In: 2015 IEEE\/RSJ International Conference on Intelligent Robots and Systems (IROS). pp 922\u2013928. https:\/\/doi.org\/10.1109\/IROS.2015.7353481","DOI":"10.1109\/IROS.2015.7353481"},{"key":"18653_CR15","doi-asserted-by":"publisher","unstructured":"Klokov R, Lempitsky V (2017) Escape from cells: deep Kd-networks for the recognition of 3D point cloud models. In: 2017 IEEE International Conference on Computer Vision (ICCV). IEEE, Venice, pp 863\u2013872. https:\/\/doi.org\/10.1109\/ICCV.2017.99","DOI":"10.1109\/ICCV.2017.99"},{"key":"18653_CR16","doi-asserted-by":"publisher","unstructured":"Riegler G, Ulusoy AO, Geiger A (2017) OctNet: learning deep 3D representations at high resolutions. In: 2017 IEEE Conference on Computer Vision and Pattern Recognition (CVPR). IEEE, Honolulu, HI, pp 6620\u20136629. https:\/\/doi.org\/10.1109\/CVPR.2017.701","DOI":"10.1109\/CVPR.2017.701"},{"key":"18653_CR17","unstructured":"Qi CR, Yi L, Su H, Guibas LJ (2017) PointNet++: deep hierarchical feature learning on point sets in a metric space. In: Advances in Neural Information Processing Systems. Curran Associates, Inc."},{"key":"18653_CR18","doi-asserted-by":"publisher","unstructured":"Zhao H, Jiang L, Fu C-W, Jia J (2019) PointWeb: enhancing local neighborhood features for point cloud processing. In: 2019 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR). IEEE, Long Beach, CA, USA, pp 5560\u20135568. https:\/\/doi.org\/10.1109\/CVPR.2019.00571","DOI":"10.1109\/CVPR.2019.00571"},{"key":"18653_CR19","doi-asserted-by":"publisher","unstructured":"Wang H, Liu Q, Yue X et al (2021) Unsupervised point cloud pre-training via occlusion completion. In: 2021 IEEE\/CVF International Conference on Computer Vision (ICCV). IEEE, Montreal, QC, Canada, pp 9762\u20139772. https:\/\/doi.org\/10.1109\/ICCV48922.2021.00964","DOI":"10.1109\/ICCV48922.2021.00964"},{"key":"18653_CR20","doi-asserted-by":"publisher","unstructured":"Poursaeed O, Jiang T, Qiao H et al (2020) Self-Supervised learning of point clouds via orientation estimation. In: 2020 International Conference on 3D Vision (3DV). pp 1018\u20131028. https:\/\/doi.org\/10.1109\/3DV50981.2020.00112","DOI":"10.1109\/3DV50981.2020.00112"},{"key":"18653_CR21","doi-asserted-by":"publisher","unstructured":"He K, Fan H, Wu Y et al (2020) Momentum contrast for unsupervised visual representation learning. In: 2020 IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR). IEEE, Seattle, WA, USA, pp 9726\u20139735. https:\/\/doi.org\/10.1109\/CVPR42600.2020.00975","DOI":"10.1109\/CVPR42600.2020.00975"},{"key":"18653_CR22","unstructured":"Chen T, Kornblith S, Norouzi M, Hinton G (2020) A simple framework for contrastive learning of visual representations. In: Proceedings of the 37th International Conference on Machine Learning. PMLR, pp 1597\u20131607"},{"key":"18653_CR23","doi-asserted-by":"publisher","unstructured":"Oord AVD, Li Y, Vinyals O (2018) Representation learning with contrastive predictive coding. arXiv preprint arXiv:1807.03748. https:\/\/doi.org\/10.48550\/arXiv.1807.03748","DOI":"10.48550\/arXiv.1807.03748"},{"key":"18653_CR24","doi-asserted-by":"publisher","first-page":"5","DOI":"10.1007\/BF03024331","volume":"19","author":"EB Saff","year":"1997","unstructured":"Saff EB, Kuijlaars ABJ (1997) Distributing many points on a sphere. Math Intelligencer 19:5\u201311. https:\/\/doi.org\/10.1007\/BF03024331","journal-title":"Math Intelligencer"},{"key":"18653_CR25","doi-asserted-by":"publisher","first-page":"49","DOI":"10.1007\/s11004-009-9257-x","volume":"42","author":"\u00c1 Gonz\u00e1lez","year":"2010","unstructured":"Gonz\u00e1lez \u00c1 (2010) Measurement of areas on a sphere using Fibonacci and latitude\u2013longitude lattices. Math Geosci 42:49\u201364. https:\/\/doi.org\/10.1007\/s11004-009-9257-x","journal-title":"Math Geosci"},{"key":"18653_CR26","doi-asserted-by":"publisher","unstructured":"Lazzarotto D, Ebrahimi T (2022) Sampling color and geometry point clouds from ShapeNet dataset. arXiv preprint arXiv:2201.06935. https:\/\/doi.org\/10.48550\/arXiv.2201.06935","DOI":"10.48550\/arXiv.2201.06935"},{"key":"18653_CR27","doi-asserted-by":"publisher","unstructured":"Uy MA, Pham Q-H, Hua B-S et al (2019) Revisiting point cloud classification: a new benchmark dataset and classification model on real-world data. In: 2019 IEEE\/CVF International Conference on Computer Vision (ICCV). IEEE, Seoul, Korea (South), pp 1588\u20131597. https:\/\/doi.org\/10.1109\/ICCV.2019.00167","DOI":"10.1109\/ICCV.2019.00167"},{"key":"18653_CR28","doi-asserted-by":"publisher","unstructured":"Hua B-S, Pham Q-H, Nguyen DT et al (2016) SceneNN: a scene meshes dataset with aNNotations. In: 2016 Fourth International Conference on 3D Vision (3DV). pp 92\u2013101. https:\/\/doi.org\/10.1109\/3DV.2016.18","DOI":"10.1109\/3DV.2016.18"},{"key":"18653_CR29","doi-asserted-by":"publisher","unstructured":"Dai A, Chang AX, Savva M, et al (2017) ScanNet: richly-annotated 3D reconstructions of indoor scenes. In: 2017 IEEE Conference on Computer Vision and Pattern Recognition (CVPR). IEEE, Honolulu, HI, pp 2432\u20132443. https:\/\/doi.org\/10.1109\/CVPR.2017.261","DOI":"10.1109\/CVPR.2017.261"},{"key":"18653_CR30","doi-asserted-by":"publisher","unstructured":"Goyal P, Doll\u00e1r P, Girshick R, Noordhuis P, Wesolowski L, Kyrola A, Tulloch A, Jia Y, He K (2017) Accurate, large minibatch sgd: training imagenet in 1 hour. arXiv preprint arXiv:1706.02677. https:\/\/doi.org\/10.48550\/arXiv.1706.02677","DOI":"10.48550\/arXiv.1706.02677"},{"key":"18653_CR31","doi-asserted-by":"publisher","unstructured":"Johnson J, Ravi N, Reizenstein J, Novotny D, Tulsiani S, Lassner C, Branson S (2020) Accelerating 3d deep learning with pytorch3d. In: SIGGRAPH Asia 2020 Courses. pp 1\u20131. https:\/\/doi.org\/10.1145\/3415263.3419160","DOI":"10.1145\/3415263.3419160"},{"key":"18653_CR32","doi-asserted-by":"publisher","unstructured":"Hassani K, Haley M (2019) Unsupervised multi-task feature learning on point clouds. In: 2019 IEEE\/CVF International Conference on Computer Vision (ICCV). IEEE, Seoul, Korea (South), pp 8159\u20138170. https:\/\/doi.org\/10.1109\/ICCV.2019.00825","DOI":"10.1109\/ICCV.2019.00825"},{"key":"18653_CR33","unstructured":"Sauder J, Sievers B (2019) Self-supervised deep learning on point clouds by reconstructing space. Adv Neural Inf Proces Syst 32"},{"key":"18653_CR34","unstructured":"Sharma C, Kaul M (2020) Self-supervised few-shot learning on point clouds. Adv Neural Inf Proces Systs 33:7212\u20137221"}],"container-title":["Multimedia Tools and Applications"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11042-024-18653-7.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11042-024-18653-7\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11042-024-18653-7.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2024,9,4]],"date-time":"2024-09-04T04:25:34Z","timestamp":1725423934000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11042-024-18653-7"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2024,2,27]]},"references-count":34,"journal-issue":{"issue":"32","published-online":{"date-parts":[[2024,9]]}},"alternative-id":["18653"],"URL":"https:\/\/doi.org\/10.1007\/s11042-024-18653-7","relation":{},"ISSN":["1573-7721"],"issn-type":[{"value":"1573-7721","type":"electronic"}],"subject":[],"published":{"date-parts":[[2024,2,27]]},"assertion":[{"value":"13 June 2023","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"13 February 2024","order":2,"name":"revised","label":"Revised","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"14 February 2024","order":3,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"27 February 2024","order":4,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors have no competing interests to declare that are relevant to the content of this article.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}]}}