{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,9]],"date-time":"2026-07-09T20:19:20Z","timestamp":1783628360606,"version":"3.55.0"},"reference-count":62,"publisher":"Springer Science and Business Media LLC","issue":"12","license":[{"start":{"date-parts":[[2025,11,18]],"date-time":"2025-11-18T00:00:00Z","timestamp":1763424000000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2025,11,18]],"date-time":"2025-11-18T00:00:00Z","timestamp":1763424000000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/100009044","name":"Technische Universit\u00e4t Kaiserslautern","doi-asserted-by":"publisher","id":[{"id":"10.13039\/100009044","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Mach Learn"],"published-print":{"date-parts":[[2025,12]]},"DOI":"10.1007\/s10994-025-06903-0","type":"journal-article","created":{"date-parts":[[2025,11,18]],"date-time":"2025-11-18T23:13:13Z","timestamp":1763507593000},"update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":4,"title":["Multi-modal co-learning for Earth observation: enhancing single-modality models via modality collaboration"],"prefix":"10.1007","volume":"114","author":[{"ORCID":"https:\/\/orcid.org\/0000-0002-5004-6571","authenticated-orcid":false,"given":"Francisco","family":"Mena","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-8736-3132","authenticated-orcid":false,"given":"Dino","family":"Ienco","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-1934-0625","authenticated-orcid":false,"given":"C\u00e0ssio F.","family":"Dantas","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-0536-6277","authenticated-orcid":false,"given":"Roberto","family":"Interdonato","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-6100-8255","authenticated-orcid":false,"given":"Andreas","family":"Dengel","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2025,11,18]]},"reference":[{"key":"6903_CR1","unstructured":"Andrew, G., Arora, R., Bilmes, J., & Livescu, K. (2013). Deep canonical correlation analysis. In International Conference on Machine Learning (pp. 1247\u20131255)."},{"key":"6903_CR2","doi-asserted-by":"crossref","unstructured":"Astruc, G., Gonthier, N., & Mallet, C., Landrieu, L. (2025). AnySat: An Earth observation model for any resolutions, scales, and modalities. In IEEE\/CVF Conference on Computer Vision and Pattern Recognition.","DOI":"10.1109\/CVPR52734.2025.01819"},{"key":"6903_CR3","doi-asserted-by":"publisher","unstructured":"Astruc, G., Gonthier, N., Mallet, C., & Landrieu, L. (2025). OmniSat: Self-supervised modality fusion for Earth observation. In Proceedings of the European Conference on Computer Vision (pp. 409\u2013427). https:\/\/doi.org\/10.1007\/978-3-031-73390-1_24","DOI":"10.1007\/978-3-031-73390-1_24"},{"key":"6903_CR4","doi-asserted-by":"publisher","unstructured":"Bakalos, N., Sykiotis, S., Temenos, A., Rallis, I., Doulamis, A., & Doulamis, N. (2024). Segmentation of remote sensing data with missing modalities through prototype knowledge distillation. In IEEE International Geoscience and Remote Sensing Symposium (pp. 10015\u201310018). https:\/\/doi.org\/10.1109\/IGARSS53475.2024.10641049","DOI":"10.1109\/IGARSS53475.2024.10641049"},{"issue":"2","key":"6903_CR5","doi-asserted-by":"publisher","first-page":"423","DOI":"10.1109\/TPAMI.2018.2798607","volume":"41","author":"T Baltru\u0161aitis","year":"2018","unstructured":"Baltru\u0161aitis, T., Ahuja, C., & Morency, L.-P. (2018). Multimodal machine learning: a survey and taxonomy. IEEE Transactions on Pattern Analysis and Machine Intelligence, 41(2), 423\u2013443. https:\/\/doi.org\/10.1109\/TPAMI.2018.2798607","journal-title":"IEEE Transactions on Pattern Analysis and Machine Intelligence"},{"key":"6903_CR6","doi-asserted-by":"publisher","unstructured":"Black, S., & Souvenir, R. (2024). Multi-view classification using hybrid fusion and mutual distillation. In Proceedings of the IEEE\/CVF Winter Conference on Applications of Computer Vision (pp. 270\u2013280). https:\/\/doi.org\/10.1109\/WACV57701.2024.00034","DOI":"10.1109\/WACV57701.2024.00034"},{"key":"6903_CR7","doi-asserted-by":"crossref","unstructured":"Blum, A., & Mitchell, T. (1998). Combining labeled and unlabeled data with co-training. In Proceedings of the Eleventh Annual Conference on Computational Learning Theory (pp. 92\u2013100).","DOI":"10.1145\/279943.279962"},{"key":"6903_CR8","series-title":"Climate science and geosciences","doi-asserted-by":"publisher","DOI":"10.1002\/9781119646181","volume-title":"Deep learning for the earth sciences: A comprehensive approach to remote sensing","author":"G Camps-Valls","year":"2021","unstructured":"Camps-Valls, G., Tuia, D., Zhu, X. X., & Reichstein, M. (2021). Deep learning for the earth sciences: A comprehensive approach to remote sensing. Climate science and geosciencesHoboken: Wiley."},{"key":"6903_CR9","doi-asserted-by":"publisher","unstructured":"Carrascosa, C., Rinc\u00f3n, J., & Rebollo, M. (2022). Co-learning: Consensus-based learning for multi-agent systems. In International Conference on Practical Applications of Agents and Multi-Agent Systems (pp. 63\u201375). https:\/\/doi.org\/10.1007\/978-3-031-18192-4_6","DOI":"10.1007\/978-3-031-18192-4_6"},{"key":"6903_CR10","unstructured":"Chen, T., Kornblith, S., Norouzi, M., & Hinton, G. (2020). A simple framework for contrastive learning of visual representations. In International Conference on Machine Learning (pp. 1597\u20131607)."},{"issue":"10","key":"6903_CR11","doi-asserted-by":"publisher","first-page":"7049","DOI":"10.1109\/TGRS.2020.2979273","volume":"58","author":"Y Chen","year":"2020","unstructured":"Chen, Y., Lu, X., & Wang, S. (2020). Deep cross-modal image-voice retrieval in remote sensing. IEEE Transactions on Geoscience and Remote Sensing, 58(10), 7049\u20137061. https:\/\/doi.org\/10.1109\/TGRS.2020.2979273","journal-title":"IEEE Transactions on Geoscience and Remote Sensing"},{"key":"6903_CR12","doi-asserted-by":"publisher","first-page":"5404914","DOI":"10.1109\/TGRS.2024.3387837","volume":"62","author":"Y Chen","year":"2024","unstructured":"Chen, Y., Zhao, M., & Bruzzone, L. (2024). A novel approach to incomplete multimodal learning for remote sensing data fusion. IEEE Transactions on Geoscience and Remote Sensing, 62, 5404914. https:\/\/doi.org\/10.1109\/TGRS.2024.3387837","journal-title":"IEEE Transactions on Geoscience and Remote Sensing"},{"key":"6903_CR13","doi-asserted-by":"publisher","first-page":"259","DOI":"10.1016\/j.inffus.2019.02.010","volume":"51","author":"J-H Choi","year":"2019","unstructured":"Choi, J.-H., & Lee, J.-S. (2019). EmbraceNet: A robust deep learning architecture for multimodal classification. Information Fusion, 51, 259\u2013270. https:\/\/doi.org\/10.1016\/j.inffus.2019.02.010","journal-title":"Information Fusion"},{"key":"6903_CR14","doi-asserted-by":"publisher","first-page":"1681","DOI":"10.1109\/JSTARS.2024.3503756","volume":"18","author":"CF Dantas","year":"2024","unstructured":"Dantas, C. F., Gaetano, R., Paris, C., & Ienco, D. (2024). Reuse out-of-year data to enhance land cover mapping via feature disentanglement and contrastive learning. IEEE Journal of Selected Topics in Applied Earth Observations and Remote Sensing, 18, 1681\u20131694. https:\/\/doi.org\/10.1109\/JSTARS.2024.3503756","journal-title":"IEEE Journal of Selected Topics in Applied Earth Observations and Remote Sensing"},{"issue":"1","key":"6903_CR15","doi-asserted-by":"publisher","first-page":"22","DOI":"10.1007\/s10618-023-00948-2","volume":"38","author":"NM Foumani","year":"2024","unstructured":"Foumani, N. M., Tan, C. W., Webb, G. I., & Salehi, M. (2024). Improving position encoding of transformers for multivariate time series classification. Data Mining and Knowledge Discovery, 38(1), 22\u201348. https:\/\/doi.org\/10.1007\/s10618-023-00948-2","journal-title":"Data Mining and Knowledge Discovery"},{"key":"6903_CR16","unstructured":"Ganin, Y., & Lempitsky, V. (2015). Unsupervised domain adaptation by backpropagation. In International Conference on Machine Learning (pp. 1180\u20131189)."},{"key":"6903_CR17","unstructured":"Han, B., Yao, Q., Yu, X., Niu, G., Xu, M., Hu, W., Tsang, I., & Sugiyama, M. (2018) Co-teaching: Robust training of deep neural networks with extremely noisy labels. Advances in Neural Information Processing Systems 31."},{"key":"6903_CR18","doi-asserted-by":"publisher","unstructured":"Hazarika, D., Zimmermann, R., & Poria, S. (2020). MISA: Modality-invariant and-specific representations for multimodal sentiment analysis. In Proceedings of the 28th ACM International Conference on Multimedia (pp. 1122\u20131131). https:\/\/doi.org\/10.1145\/3394171.341367","DOI":"10.1145\/3394171.341367"},{"key":"6903_CR19","doi-asserted-by":"publisher","unstructured":"He, K., Zhang, X., Ren, S., & Sun, J. (2016). Deep residual learning for image recognition. In Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (pp. 770\u2013778). https:\/\/doi.org\/10.1109\/CVPR.2016.90","DOI":"10.1109\/CVPR.2016.90"},{"key":"6903_CR20","doi-asserted-by":"publisher","DOI":"10.1016\/j.jag.2022.103130","volume":"116","author":"K Heidler","year":"2023","unstructured":"Heidler, K., Mou, L., Hu, D., Jin, P., Li, G., Gan, C., Wen, J.-R., & Zhu, X. X. (2023). Self-supervised audiovisual representation learning for remote sensing data. International Journal of Applied Earth Observation and Geoinformation, 116, Article 103130. https:\/\/doi.org\/10.1016\/j.jag.2022.103130","journal-title":"International Journal of Applied Earth Observation and Geoinformation"},{"issue":"8","key":"6903_CR21","doi-asserted-by":"publisher","first-page":"1735","DOI":"10.1162\/neco.1997.9.8.1735","volume":"9","author":"S Hochreiter","year":"1997","unstructured":"Hochreiter, S., & Schmidhuber, J. (1997). Long short-term memory. Neural Computation, 9(8), 1735\u20131780. https:\/\/doi.org\/10.1162\/neco.1997.9.8.1735","journal-title":"Neural Computation"},{"issue":"7","key":"6903_CR22","doi-asserted-by":"publisher","first-page":"4349","DOI":"10.1109\/TGRS.2018.2890705","volume":"57","author":"D Hong","year":"2019","unstructured":"Hong, D., Yokoya, N., Chanussot, J., & Zhu, X. X. (2019). Cospace: Common subspace learning from hyperspectral-multispectral correspondences. IEEE Transactions on Geoscience and Remote Sensing, 57(7), 4349\u20134359. https:\/\/doi.org\/10.1109\/TGRS.2018.2890705","journal-title":"IEEE Transactions on Geoscience and Remote Sensing"},{"key":"6903_CR23","unstructured":"Ienco, D., & Dantas, C.F. (2024). DisCoM-KD: Cross-modal knowledge distillation via disentanglement representation and adversarial learning. In British Machine Vision Conference."},{"key":"6903_CR24","unstructured":"Jiang, L., Zhou, Z., Leung, T., Li, L.-J., & Fei-Fei, L. (2018). MentorNet: Learning data-driven curriculum for very deep neural networks on corrupted labels. In International Conference on Machine Learning (pp. 2304\u20132313)."},{"issue":"6","key":"6903_CR25","doi-asserted-by":"publisher","first-page":"1758","DOI":"10.1109\/JSTARS.2018.2834961","volume":"11","author":"M Kampffmeyer","year":"2018","unstructured":"Kampffmeyer, M., Salberg, A.-B., & Jenssen, R. (2018). Urban land cover classification with missing data modalities using deep convolutional neural networks. IEEE Journal of Selected Topics in Applied Earth Observations and Remote Sensing, 11(6), 1758\u20131768. https:\/\/doi.org\/10.1109\/JSTARS.2018.2834961","journal-title":"IEEE Journal of Selected Topics in Applied Earth Observations and Remote Sensing"},{"key":"6903_CR26","doi-asserted-by":"publisher","unstructured":"Kendall, A., Gal, Y., & Cipolla, R. (2018). Multi-task learning using uncertainty to weigh losses for scene geometry and semantics. In Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (pp. 7482\u20137491). https:\/\/doi.org\/10.1109\/CVPR.2018.00781","DOI":"10.1109\/CVPR.2018.00781"},{"key":"6903_CR27","doi-asserted-by":"publisher","DOI":"10.1109\/JSTARS.2024.3378348","author":"N Kieu","year":"2024","unstructured":"Kieu, N., Nguyen, K., Nazib, A., Fernando, T., Fookes, C., & Sridharan, S. (2024). Multimodal co-learning meets remote sensing: Taxonomy, state of the art, and future works. IEEE Journal of Selected Topics in Applied Earth Observations and Remote Sensing. https:\/\/doi.org\/10.1109\/JSTARS.2024.3378348","journal-title":"IEEE Journal of Selected Topics in Applied Earth Observations and Remote Sensing"},{"key":"6903_CR28","doi-asserted-by":"publisher","first-page":"1686","DOI":"10.1162\/tacl_a_00628","volume":"11","author":"R Lin","year":"2023","unstructured":"Lin, R., & Hu, H. (2023). MissModal: Increasing robustness to missing modality in multimodal sentiment analysis. Transactions of the Association for Computational Linguistics, 11, 1686\u20131702. https:\/\/doi.org\/10.1162\/tacl_a_00628","journal-title":"Transactions of the Association for Computational Linguistics"},{"issue":"11","key":"6903_CR29","first-page":"2579","volume":"9","author":"L Maaten","year":"2008","unstructured":"Maaten, L., & Hinton, G. (2008). Visualizing data using t-SNE. Journal of Machine Learning Research, 9(11), 2579\u20132605.","journal-title":"Journal of Machine Learning Research"},{"issue":"12","key":"6903_CR30","doi-asserted-by":"publisher","first-page":"2691","DOI":"10.1109\/TGRS.2004.840720","volume":"42","author":"BL Markham","year":"2004","unstructured":"Markham, B. L., Storey, J. C., Williams, D. L., & Irons, J. R. (2004). Landsat sensor performance: History and current status. IEEE Transactions on Geoscience and Remote Sensing, 42(12), 2691\u20132694. https:\/\/doi.org\/10.1109\/TGRS.2004.840720","journal-title":"IEEE Transactions on Geoscience and Remote Sensing"},{"key":"6903_CR31","unstructured":"McKinzie, B., Shankar, V., Cheng, J.Y., Yang, Y., Shlens, J., & Toshev, A.T. (2023). Robustness in multimodal learning under train-test modality mismatch. In International Conference on Machine Learning (pp. 24291\u201324303)."},{"key":"6903_CR32","unstructured":"Mena, F., Arenas, D., & Dengel, A. (2024). Increasing the robustness of model predictions to missing sensors in Earth observation. arXiv preprint arXiv:2407.15512"},{"key":"6903_CR33","doi-asserted-by":"publisher","DOI":"10.1016\/j.neucom.2025.130175","volume":"638","author":"F Mena","year":"2025","unstructured":"Mena, F., Arenas, D., & Dengel, A. (2025). Missing data as augmentation in the Earth observation domain: A multi-view learning approach. Neurocomputing, 638, Article 130175. https:\/\/doi.org\/10.1016\/j.neucom.2025.130175","journal-title":"Neurocomputing"},{"key":"6903_CR34","doi-asserted-by":"publisher","first-page":"4797","DOI":"10.1109\/JSTARS.2024.3361556","volume":"17","author":"F Mena","year":"2024","unstructured":"Mena, F., Arenas, D., Nuske, M., & Dengel, A. (2024). Common practices and taxonomy in deep multi-view fusion for remote sensing applications. IEEE Journal of Selected Topics in Applied Earth Observations and Remote Sensing, 17, 4797\u20134818. https:\/\/doi.org\/10.1109\/JSTARS.2024.3361556","journal-title":"IEEE Journal of Selected Topics in Applied Earth Observations and Remote Sensing"},{"key":"6903_CR35","doi-asserted-by":"publisher","DOI":"10.1016\/j.rse.2024.114547","volume":"318","author":"F Mena","year":"2025","unstructured":"Mena, F., Pathak, D., Najjar, H., Sanchez, C., Helber, P., Bischke, B., Habelitz, P., Miranda, M., Siddamsetty, J., Nuske, M., Charfuelan, M., Arenas, D., Vollmer, M., & Dengel, A. (2025). Adaptive fusion of multi-modal remote sensing data for optimal sub-field crop yield prediction. Remote Sensing of Environment, 318, Article 114547. https:\/\/doi.org\/10.1016\/j.rse.2024.114547","journal-title":"Remote Sensing of Environment"},{"key":"6903_CR36","doi-asserted-by":"publisher","unstructured":"Nedungadi, V., Kariryaa, A., Oehmcke, S., Belongie, S., Igel, C., & Lang, N. (2024). MMEarth: Exploring multi-modal pretext tasks for geospatial representation learning. In Proceedings of the European Conference on Computer Vision (pp. 164\u2013182). https:\/\/doi.org\/10.1007\/978-3-031-73039-9_10","DOI":"10.1007\/978-3-031-73039-9_10"},{"key":"6903_CR37","unstructured":"Ngiam, J., Khosla, A., Kim, M., Nam, J., Lee, H., & Ng, A.Y. (2011). Multimodal deep learning. In International Conference on Machine Learning vol. 11 (pp. 689\u2013696)."},{"issue":"10","key":"6903_CR38","doi-asserted-by":"publisher","first-page":"3659","DOI":"10.1007\/s10994-023-06374-1","volume":"112","author":"M Obrenovi\u0107","year":"2023","unstructured":"Obrenovi\u0107, M., Lampert, T., Ivanovi\u0107, M., & Gan\u00e7arski, P. (2023). Learning domain invariant representations of heterogeneous image data. Machine Learning, 112(10), 3659\u20133684. https:\/\/doi.org\/10.1007\/s10994-023-06374-1","journal-title":"Machine Learning"},{"key":"6903_CR39","doi-asserted-by":"publisher","unstructured":"Pande, S., Banerjee, A., Kumar, S., Banerjee, B., & Chaudhuri, S. (2019). An adversarial approach to discriminative modality distillation for remote sensing image classification. In Proceedings of the IEEE\/CVF International Conference on Computer Vision Workshops. https:\/\/doi.org\/10.1109\/ICCVW.2019.00558","DOI":"10.1109\/ICCVW.2019.00558"},{"issue":"5","key":"6903_CR40","doi-asserted-by":"publisher","first-page":"523","DOI":"10.3390\/rs11050523","volume":"11","author":"C Pelletier","year":"2019","unstructured":"Pelletier, C., Webb, G. I., & Petitjean, F. (2019). Temporal convolutional neural network for the classification of satellite image time series. Remote Sensing, 11(5), 523. https:\/\/doi.org\/10.3390\/rs11050523","journal-title":"Remote Sensing"},{"issue":"1","key":"6903_CR41","doi-asserted-by":"publisher","first-page":"94","DOI":"10.1109\/MGRS.2024.3367006","volume":"11","author":"C Persello","year":"2023","unstructured":"Persello, C., H\u00e4nsch, R., Vivone, G., Chen, K., Yan, Z., Tang, D., Huang, H., Schmitt, M., & Sun, X. (2023). 2023 IEEE GRSS data fusion contest: Large-scale fine-grained building classification for semantic urban reconstruction. IEEE Geoscience and Remote Sensing Magazine, 11(1), 94\u201397. https:\/\/doi.org\/10.1109\/MGRS.2024.3367006","journal-title":"IEEE Geoscience and Remote Sensing Magazine"},{"issue":"2","key":"6903_CR42","doi-asserted-by":"publisher","first-page":"109","DOI":"10.1109\/MGRS.2023.3240233","volume":"12","author":"C Persello","year":"2024","unstructured":"Persello, C., Prasad, S., Vivone, G., Lonjou, V., Bretar, F., Rodriguez-Suquet, R., Guntzburger, P., Poulain, V., Moigne, J. L., Smith, B., et al. (2024). 2024 IEEE GRSS data fusion contest: Rapid flood mapping. IEEE Geoscience and Remote Sensing Magazine, 12(2), 109\u2013112. https:\/\/doi.org\/10.1109\/MGRS.2023.3240233","journal-title":"IEEE Geoscience and Remote Sensing Magazine"},{"key":"6903_CR43","unstructured":"Poklukar, P., Vasco, M., Yin, H., Melo, F.S., Paiva, A., & Kragic, D. (2022). Geometric multimodal contrastive representation learning. In International Conference on Machine Learning (pp. 17782\u201317800)."},{"key":"6903_CR44","doi-asserted-by":"publisher","unstructured":"Potin, P., Colin, O., Pinheiro, M., Rosich, B., O\u2019Connell, A., Ormston, T., Gratadour, J.-B., & Torres, R. (2022). Status and evolution of the Sentinel-1 mission. In IEEE International Geoscience and Remote Sensing Symposium (pp. 4707\u20134710). https:\/\/doi.org\/10.1109\/IGARSS46834.2022.9884753","DOI":"10.1109\/IGARSS46834.2022.9884753"},{"key":"6903_CR45","doi-asserted-by":"publisher","unstructured":"Qiao, S., Shen, W., Zhang, Z., Wang, B., & Yuille, A. (2018). Deep co-training for semi-supervised image recognition. In Proceedings of the European Conference on Computer Vision (pp. 135\u2013152). https:\/\/doi.org\/10.1007\/978-3-030-01267-0_9","DOI":"10.1007\/978-3-030-01267-0_9"},{"key":"6903_CR46","doi-asserted-by":"publisher","first-page":"203","DOI":"10.1016\/j.inffus.2021.12.003","volume":"81","author":"A Rahate","year":"2022","unstructured":"Rahate, A., Walambe, R., Ramanna, S., & Kotecha, K. (2022). Multimodal co-learning: challenges, applications with datasets, recent advances and future directions. Information Fusion, 81, 203\u2013239.","journal-title":"Information Fusion"},{"key":"6903_CR47","doi-asserted-by":"publisher","DOI":"10.1016\/j.rse.2020.111797","volume":"245","author":"K Rao","year":"2020","unstructured":"Rao, K., Williams, A. P., Flefil, J. F., & Konings, A. G. (2020). SAR-enhanced mapping of live fuel moisture content. Remote Sensing of Environment, 245, Article 111797. https:\/\/doi.org\/10.1016\/j.rse.2020.111797","journal-title":"Remote Sensing of Environment"},{"key":"6903_CR48","unstructured":"Rolf, E., Klemmer, K., Robinson, C., & Kerner, H. (2024). Mission critical\u2013satellite data is a distinct modality in machine learning. arXiv preprint arXiv:2402.01444"},{"issue":"3","key":"6903_CR49","doi-asserted-by":"publisher","first-page":"61","DOI":"10.1109\/MGRS.2015.2441912","volume":"3","author":"H Shen","year":"2015","unstructured":"Shen, H., Li, X., Cheng, Q., Zeng, C., Yang, G., Li, H., & Zhang, L. (2015). Missing information reconstruction of remote sensing data: A technical review. IEEE Geoscience and Remote Sensing Magazine, 3(3), 61\u201385. https:\/\/doi.org\/10.1109\/MGRS.2015.2441912","journal-title":"IEEE Geoscience and Remote Sensing Magazine"},{"key":"6903_CR50","unstructured":"Tseng, G., Zvonkov, I., Nakalembe, C.L., & Kerner, H. (2021). CropHarvest: A global dataset for crop-type classification. In Proceedings of NIPS Datasets and Benchmarks Track."},{"issue":"2","key":"6903_CR51","doi-asserted-by":"publisher","first-page":"88","DOI":"10.1109\/MGRS.2020.3043504","volume":"9","author":"D Tuia","year":"2021","unstructured":"Tuia, D., Roscher, R., Wegner, J. D., Jacobs, N., Zhu, X., & Camps-Valls, G. (2021). Toward a collective agenda on ai for earth science data analysis. IEEE Geoscience and Remote Sensing Magazine, 9(2), 88\u2013104. https:\/\/doi.org\/10.1109\/MGRS.2020.3043504","journal-title":"IEEE Geoscience and Remote Sensing Magazine"},{"key":"6903_CR52","unstructured":"Vaswani, A., Shazeer, N., Parmar, N., Uszkoreit, J., Jones, L., Gomez, A.N., Kaiser, \u0141., & Polosukhin, I. (2017). Attention is all you need. In Advances in Neural Information Processing Systems 30."},{"key":"6903_CR53","doi-asserted-by":"publisher","unstructured":"Wang, Y., Albrecht, C.M., Braham, N.A.A., Liu, C., Xiong, Z., & Zhu, X.X. (2025). DeCUR: Decoupling common & unique representations for multimodal self-supervision. In Proceedings of the European Conference on Computer Vision (pp. 286\u2013303). https:\/\/doi.org\/10.1007\/978-3-031-73397-0_17","DOI":"10.1007\/978-3-031-73397-0_17"},{"key":"6903_CR54","doi-asserted-by":"publisher","unstructured":"Wang, H., Chen, Y., Ma, C., Avery, J., Hull, L., & Carneiro, G. (2023). Multi-modal learning with missing modality via shared-specific feature modelling. In Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (pp. 15878\u201315887). https:\/\/doi.org\/10.1109\/CVPR52729.2023.01524","DOI":"10.1109\/CVPR52729.2023.01524"},{"key":"6903_CR55","unstructured":"Wu, R., Wang, H., Chen, H.-T., & Carneiro, G. (2024). Deep multimodal learning with missing modality: A survey. arXiv preprint arXiv:2409.07825."},{"key":"6903_CR56","doi-asserted-by":"publisher","DOI":"10.1016\/j.jag.2022.103165","volume":"116","author":"Y Xie","year":"2023","unstructured":"Xie, Y., Tian, J., & Zhu, X. X. (2023). A co-learning method to utilize optical images and photogrammetric point clouds for building extraction. International Journal of Applied Earth Observation and Geoinformation, 116, Article 103165. https:\/\/doi.org\/10.1016\/j.jag.2022.103165","journal-title":"International Journal of Applied Earth Observation and Geoinformation"},{"key":"6903_CR57","doi-asserted-by":"publisher","unstructured":"Xiong, Z., Wang, Y., Zhang, F., & Zhu, X.X. (2024). One for all: Toward unified foundation models for Earth vision. In IEEE International Geoscience and Remote Sensing Symposium (pp. 2734\u20132738). https:\/\/doi.org\/10.1109\/IGARSS53475.2024.10641637","DOI":"10.1109\/IGARSS53475.2024.10641637"},{"key":"6903_CR58","doi-asserted-by":"publisher","unstructured":"Yuan, X., Lin, Z., Kuen, J., Zhang, J., Wang, Y., Maire, M., Kale, A., & Faieta, B. (2021). Multimodal contrastive training for visual representation learning. In Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (pp. 6995\u20137004). https:\/\/doi.org\/10.1109\/CVPR46437.2021.00692","DOI":"10.1109\/CVPR46437.2021.00692"},{"key":"6903_CR59","doi-asserted-by":"publisher","first-page":"188","DOI":"10.1016\/j.inffus.2020.06.001","volume":"64","author":"A Zadeh","year":"2020","unstructured":"Zadeh, A., Liang, P. P., & Morency, L.-P. (2020). Foundations of multimodal co-learning. Information Fusion, 64, 188\u2013193. https:\/\/doi.org\/10.1016\/j.inffus.2020.06.001","journal-title":"Information Fusion"},{"key":"6903_CR61","doi-asserted-by":"publisher","unstructured":"Zhang, X., Yoon, J., Bansal, M., & Yao, H. (2024). Multimodal representation learning by alternating unimodal adaptation. In Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (pp. 27456\u201327466). https:\/\/doi.org\/10.1109\/CVPR52733.2024.02592","DOI":"10.1109\/CVPR52733.2024.02592"},{"key":"6903_CR60","doi-asserted-by":"publisher","unstructured":"Zhang, Y., Xiang, T., Hospedales, T.M., & Lu, H. (2018). Deep mutual learning. In Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (pp. 4320\u20134328). https:\/\/doi.org\/10.1109\/CVPR.2018.00454","DOI":"10.1109\/CVPR.2018.00454"},{"key":"6903_CR62","doi-asserted-by":"publisher","first-page":"254","DOI":"10.1016\/j.isprsjprs.2020.12.009","volume":"174","author":"Z Zheng","year":"2021","unstructured":"Zheng, Z., Ma, A., Zhang, L., & Zhong, Y. (2021). Deep multisensor learning for missing-modality all-weather mapping. ISPRS Journal of Photogrammetry and Remote Sensing, 174, 254\u2013264. https:\/\/doi.org\/10.1016\/j.isprsjprs.2020.12.009","journal-title":"ISPRS Journal of Photogrammetry and Remote Sensing"}],"container-title":["Machine Learning"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10994-025-06903-0.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s10994-025-06903-0","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10994-025-06903-0.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,12,22]],"date-time":"2025-12-22T21:29:39Z","timestamp":1766438979000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s10994-025-06903-0"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2025,11,18]]},"references-count":62,"journal-issue":{"issue":"12","published-print":{"date-parts":[[2025,12]]}},"alternative-id":["6903"],"URL":"https:\/\/doi.org\/10.1007\/s10994-025-06903-0","relation":{},"ISSN":["0885-6125","1573-0565"],"issn-type":[{"value":"0885-6125","type":"print"},{"value":"1573-0565","type":"electronic"}],"subject":[],"published":{"date-parts":[[2025,11,18]]},"assertion":[{"value":"3 April 2025","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"25 July 2025","order":2,"name":"revised","label":"Revised","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"23 September 2025","order":3,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"18 November 2025","order":4,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no conflict of interest to report regarding the present study.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}],"article-number":"279"}}