{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,8,2]],"date-time":"2026-08-02T19:34:49Z","timestamp":1785699289655,"version":"3.56.0"},"reference-count":54,"publisher":"Springer Science and Business Media LLC","issue":"3","license":[{"start":{"date-parts":[[2024,9,30]],"date-time":"2024-09-30T00:00:00Z","timestamp":1727654400000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2024,9,30]],"date-time":"2024-09-30T00:00:00Z","timestamp":1727654400000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"publisher","award":["62372019"],"award-info":[{"award-number":["62372019"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Int J Comput Vis"],"published-print":{"date-parts":[[2025,3]]},"DOI":"10.1007\/s11263-024-02233-1","type":"journal-article","created":{"date-parts":[[2024,9,30]],"date-time":"2024-09-30T16:02:56Z","timestamp":1727712176000},"page":"1290-1305","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":9,"title":["From Gaze Jitter to Domain Adaptation: Generalizing Gaze Estimation by Manipulating High-Frequency Components"],"prefix":"10.1007","volume":"133","author":[{"ORCID":"https:\/\/orcid.org\/0000-0002-8460-8763","authenticated-orcid":false,"given":"Ruicong","family":"Liu","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-1199-9206","authenticated-orcid":false,"given":"Haofei","family":"Wang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0001-9064-7964","authenticated-orcid":false,"given":"Feng","family":"Lu","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"297","published-online":{"date-parts":[[2024,9,30]]},"reference":[{"key":"2233_CR1","doi-asserted-by":"crossref","unstructured":"Admoni, H., & Scassellati, B. (2017). Social eye gaze in human-robot interaction: a review. Journal of Human-Robot Interaction, 6(1), 25\u201363.","DOI":"10.5898\/JHRI.6.1.Admoni"},{"key":"2233_CR2","unstructured":"Biggio, B., Nelson, B., & Laskov, P. (2012) Poisoning attacks against support vector machines. arXiv preprint arXiv:1206.6389"},{"key":"2233_CR3","doi-asserted-by":"crossref","unstructured":"Cai, X., Zeng, J., Shan, S., & Chen, X. (2023). Source-free adaptive gaze estimation by uncertainty reduction. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 22035\u201322045","DOI":"10.1109\/CVPR52729.2023.02110"},{"key":"2233_CR4","unstructured":"Chen, T., Kornblith, S., Norouzi, M., & Hinton, G. (2020). A simple framework for contrastive learning of visual representations. In: International Conference on Machine Learning, pp. 1597\u20131607. PMLR"},{"key":"2233_CR5","doi-asserted-by":"crossref","unstructured":"Cheng, Y., & Lu, F. (2022) Gaze estimation using transformer. In: 2022 26th International Conference on Pattern Recognition (ICPR), pp. 3341\u20133347 . IEEE","DOI":"10.1109\/ICPR56361.2022.9956687"},{"key":"2233_CR6","doi-asserted-by":"crossref","unstructured":"Cheng, Y., Lu, F., & Zhang, X. (2018). Appearance-based gaze estimation via evaluation-guided asymmetric regression. In: Proceedings of the European Conference on Computer Vision (ECCV), pp. 100\u2013115","DOI":"10.1007\/978-3-030-01264-9_7"},{"key":"2233_CR7","doi-asserted-by":"publisher","first-page":"436","DOI":"10.1609\/aaai.v36i1.19921","volume":"36","author":"Y Cheng","year":"2022","unstructured":"Cheng, Y., Bao, Y., & Lu, F. (2022). Puregaze: Purifying gaze feature for generalizable gaze estimation. Proceedings of the AAAI Conference on Artificial Intelligence, 36, 436\u2013443.","journal-title":"Proceedings of the AAAI Conference on Artificial Intelligence"},{"key":"2233_CR8","doi-asserted-by":"publisher","first-page":"10623","DOI":"10.1609\/aaai.v34i07.6636","volume":"34","author":"Y Cheng","year":"2020","unstructured":"Cheng, Y., Huang, S., Wang, F., Qian, C., & Lu, F. (2020). A coarse-to-fine adaptive network for appearance-based gaze estimation. Proceedings of the AAAI Conference on Artificial Intelligence, 34, 10623\u201310630.","journal-title":"Proceedings of the AAAI Conference on Artificial Intelligence"},{"key":"2233_CR9","doi-asserted-by":"crossref","unstructured":"Cui, S., Wang, S., Zhuo, J., Su, C., Huang, Q., & Tian, Q. (2020) Gradually vanishing bridge for adversarial domain adaptation. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 12455\u201312464","DOI":"10.1109\/CVPR42600.2020.01247"},{"issue":"3","key":"2233_CR10","doi-asserted-by":"publisher","first-page":"151","DOI":"10.1007\/s10339-007-0168-9","volume":"8","author":"Y Demiris","year":"2007","unstructured":"Demiris, Y. (2007). Prediction of intent in robotics and multi-agent systems. Cognitive processing, 8(3), 151\u2013158.","journal-title":"Cognitive processing"},{"key":"2233_CR11","unstructured":"Dosovitskiy, A., Beyer, L., Kolesnikov, A., Weissenborn, D., Zhai, X., Unterthiner, T., Dehghani, M., Minderer, M., Heigold, G., Gelly, S., et al. (2020) An image is worth 16$$\\times $$16 words: Transformers for image recognition at scale. arXiv preprint arXiv:2010.11929"},{"key":"2233_CR12","doi-asserted-by":"crossref","unstructured":"Funes Mora, K.A., Monay, F., & Odobez, J.-M. (2014). Eyediap: A database for the development and evaluation of gaze estimation algorithms from rgb and rgb-d cameras. In: Proceedings of the Symposium on Eye Tracking Research and Applications, pp. 255\u2013258","DOI":"10.1145\/2578153.2578190"},{"key":"2233_CR13","unstructured":"Goodfellow, I., Pouget-Abadie, J., Mirza, M., Xu, B., Warde-Farley, D., Ozair, S., Courville, A., & Bengio, Y. (2014). Generative adversarial nets. Advances in neural information processing systems 27"},{"key":"2233_CR14","unstructured":"Goodfellow, I.J., Shlens, J., Szegedy, C. (2014). Explaining and harnessing adversarial examples. arXiv preprint arXiv:1412.6572"},{"key":"2233_CR15","doi-asserted-by":"crossref","unstructured":"Grauman, K., Betke, M., Gips, J., & Bradski, G.R. (2001). Communication via eye blinks-detection and duration analysis in real time. In: Proceedings of the 2001 IEEE Computer Society Conference on Computer Vision and Pattern Recognition. CVPR 2001, vol. 1,p. IEEE","DOI":"10.1109\/CVPR.2001.990641"},{"key":"2233_CR16","doi-asserted-by":"crossref","unstructured":"Guo, Z., Yuan, Z., Zhang, C., Chi, W., Ling, Y., & Zhang, S. (2020). Domain adaptation gaze estimation by embedding with prediction consistency. In: Proceedings of the Asian Conference on Computer Vision","DOI":"10.1007\/978-3-030-69541-5_18"},{"key":"2233_CR17","doi-asserted-by":"crossref","unstructured":"Hallinan, P. W. (1991). Recognizing human eyes. Geometric Methods in Computer Vision,1570, 214\u2013226. SPIE","DOI":"10.1117\/12.48426"},{"key":"2233_CR18","doi-asserted-by":"crossref","unstructured":"He, K., Fan, H., Wu, Y., Xie, S., & Girshick, R. (2020). Momentum contrast for unsupervised visual representation learning. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 9729\u20139738","DOI":"10.1109\/CVPR42600.2020.00975"},{"key":"2233_CR19","doi-asserted-by":"crossref","unstructured":"He, K., Zhang, X., Ren, S., & Sun, J. (2016). Deep residual learning for image recognition. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 770\u2013778","DOI":"10.1109\/CVPR.2016.90"},{"issue":"07","key":"2233_CR20","doi-asserted-by":"publisher","first-page":"1009","DOI":"10.1142\/S0218001499000562","volume":"13","author":"J Huang","year":"1999","unstructured":"Huang, J., & Wechsler, H. (1999). Eye detection using optimal wavelet packets and radial basis functions (rbfs). International Journal of Pattern Recognition and Artificial Intelligence, 13(07), 1009\u20131025.","journal-title":"International Journal of Pattern Recognition and Artificial Intelligence"},{"key":"2233_CR21","unstructured":"Kang, G., Jiang, L., Wei, Y., Yang, Y., & Hauptmann, A. G. (2020). Contrastive adaptation network for single-and multi-source domain adaptation. IEEE transactions on pattern analysis and machine intelligence"},{"key":"2233_CR22","doi-asserted-by":"crossref","unstructured":"Kellnhofer, P., Recasens, A., Stent, S., Matusik, W., & Torralba, A. (2019). Gaze360: Physically unconstrained gaze estimation in the wild. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 6912\u20136921","DOI":"10.1109\/ICCV.2019.00701"},{"issue":"6","key":"2233_CR23","doi-asserted-by":"publisher","first-page":"1745","DOI":"10.1002\/cav.1745","volume":"29","author":"K Krejtz","year":"2018","unstructured":"Krejtz, K., Duchowski, A., Zhou, H., J\u00f6rg, S., & Niedzielska, A. (2018). Perceptual evaluation of synthetic gaze jitter. Computer Animation and Virtual Worlds, 29(6), 1745.","journal-title":"Computer Animation and Virtual Worlds"},{"key":"2233_CR24","unstructured":"Lee, D.-H., et al. (2013). Pseudo-label: The simple and efficient semi-supervised learning method for deep neural networks. In: Workshop on Challenges in Representation Learning, ICML, vol. 3, p. 896"},{"key":"2233_CR25","unstructured":"Li, J., Zhou, P., Xiong, C., Hoi, S.C.: Prototypical contrastive learning of unsupervised representations. arXiv preprint arXiv:2005.04966 (2020)"},{"key":"2233_CR26","unstructured":"Liu, W., Ferstl, D., Schulter, S., Zebedin, L., Fua, P., & Leistner, C. (2021) Domain adaptation for semantic segmentation via patch-wise contrastive learning. arXiv preprint arXiv:2104.11056"},{"key":"2233_CR27","doi-asserted-by":"crossref","unstructured":"Liu, Y., Liu, R., Wang, H., & Lu, F. (2021). Generalizing gaze estimation with outlier-guided collaborative adaptation. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 3835\u20133844","DOI":"10.1109\/ICCV48922.2021.00381"},{"key":"2233_CR28","doi-asserted-by":"publisher","first-page":"5769","DOI":"10.1109\/TIP.2021.3082317","volume":"30","author":"A Liu","year":"2021","unstructured":"Liu, A., Liu, X., Yu, H., Zhang, C., Liu, Q., & Tao, D. (2021). Training robust deep neural networks via adversarial noise propagation. IEEE Transactions on Image Processing, 30, 5769\u20135781.","journal-title":"IEEE Transactions on Image Processing"},{"key":"2233_CR29","doi-asserted-by":"publisher","first-page":"3621","DOI":"10.1609\/aaai.v38i4.28151","volume":"38","author":"H Liu","year":"2024","unstructured":"Liu, H., Qi, J., Li, Z., Hassanpour, M., Wang, Y., Plataniotis, K. N., & Yu, Y. (2024). Test-time personalization with meta prompt for gaze estimation. Proceedings of the AAAI Conference on Artificial Intelligence, 38, 3621\u20133629.","journal-title":"Proceedings of the AAAI Conference on Artificial Intelligence"},{"key":"2233_CR30","unstructured":"Madry, A., Makelov, A., Schmidt, L., Tsipras, D., & Vladu, A. (2017) Towards deep learning models resistant to adversarial attacks. arXiv preprint arXiv:1706.06083"},{"key":"2233_CR31","unstructured":"Madry, A., Makelov, A., Schmidt, L., Tsipras, D., Vladu, A. (2017). Towards deep learning models resistant to adversarial attacks. arXiv preprint arXiv:1706.06083"},{"key":"2233_CR32","doi-asserted-by":"publisher","DOI":"10.1007\/978-1-4471-6392-3_3","volume-title":"Advances in Physiological Computing Human-Computer Interaction Series","author":"P Majaranta","year":"2014","unstructured":"Majaranta, P., & Bulling, A. (2014). Eye Tracking and Eye-Based Human-Computer Interaction. In S. Fairclough & K. Gilleade (Eds.), Advances in Physiological Computing Human-Computer Interaction Series. London: Springer. https:\/\/doi.org\/10.1007\/978-1-4471-6392-3_3"},{"key":"2233_CR33","doi-asserted-by":"publisher","DOI":"10.1016\/j.patcog.2020.107332","volume":"110","author":"X Ma","year":"2021","unstructured":"Ma, X., Niu, Y., Gu, L., Wang, Y., Zhao, Y., Bailey, J., & Lu, F. (2021). Understanding adversarial attacks on deep learning based medical image analysis systems. Pattern Recognition, 110, 107332.","journal-title":"Pattern Recognition"},{"key":"2233_CR34","doi-asserted-by":"crossref","unstructured":"Mei, S., & Zhu, X. (2015). Using machine teaching to identify optimal training-set attacks on machine learners. In: Twenty-Ninth AAAI Conference on Artificial Intelligence","DOI":"10.1609\/aaai.v29i1.9569"},{"key":"2233_CR35","doi-asserted-by":"crossref","unstructured":"Moosavi-Dezfooli, S.-M., Fawzi, A., & Frossard, P. (2016). Deepfool: a simple and accurate method to fool deep neural networks. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 2574\u20132582","DOI":"10.1109\/CVPR.2016.282"},{"key":"2233_CR36","doi-asserted-by":"crossref","unstructured":"Mughrabi, M.H., Mutasim, A.K., Stuerzlinger, W., & Batmaz, A.U. (2022). My eyes hurt: Effects of jitter in 3d gaze tracking. In: 2022 IEEE Conference on Virtual Reality and 3D User Interfaces Abstracts and Workshops (VRW), pp. 310\u2013315. IEEE","DOI":"10.1109\/VRW55335.2022.00070"},{"key":"2233_CR37","doi-asserted-by":"crossref","unstructured":"Park, H.S., Jain, E., Sheikh, Y. (2013). Predicting primary gaze behavior using social saliency fields. In: Proceedings of the IEEE International Conference on Computer Vision, pp. 3503\u20133510","DOI":"10.1109\/ICCV.2013.435"},{"key":"2233_CR38","doi-asserted-by":"crossref","unstructured":"Park, S., Mello, S. D., Molchanov, P., Iqbal, U., Hilliges, O., & Kautz, J. (2019). Few-shot adaptive gaze estimation. In: Proceedings of the IEEE\/CVF International Conference on Computer Vision, pp. 9368\u20139377","DOI":"10.1109\/ICCV.2019.00946"},{"key":"2233_CR39","doi-asserted-by":"crossref","unstructured":"Schroff, F., Kalenichenko, D., & Philbin, J. (2015). Facenet: A unified embedding for face recognition and clustering. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 815\u2013823","DOI":"10.1109\/CVPR.2015.7298682"},{"key":"2233_CR40","unstructured":"Szegedy, C., Zaremba, W., Sutskever, I., Bruna, J., Erhan, D., Goodfellow, I., & Fergus, R. (2013). Intriguing properties of neural networks. arXiv preprint arXiv:1312.6199"},{"key":"2233_CR41","unstructured":"Tarvainen, A., & Valpola, H. (2017). Mean teachers are better role models: Weight-averaged consistency targets improve semi-supervised deep learning results. Advances in neural information processing systems 30"},{"key":"2233_CR42","doi-asserted-by":"crossref","unstructured":"Terzio\u011flu, Y., Mutlu, B., & \u015eahin, E. (2020). Designing social cues for collaborative robots: the role of gaze and breathing in human-robot collaboration. In: Proceedings of the 2020 ACM\/IEEE International Conference on Human-Robot Interaction, pp. 343\u2013357","DOI":"10.1145\/3319502.3374829"},{"key":"2233_CR43","doi-asserted-by":"crossref","unstructured":"Tzeng, E., Hoffman, J., Saenko, K., & Darrell, T. (2017). Adversarial discriminative domain adaptation. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 7167\u20137176","DOI":"10.1109\/CVPR.2017.316"},{"key":"2233_CR44","doi-asserted-by":"crossref","unstructured":"Wang, X., & He, K. (2021). Enhancing the transferability of adversarial attacks through variance tuning. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 1924\u20131933","DOI":"10.1109\/CVPR46437.2021.00196"},{"key":"2233_CR45","doi-asserted-by":"crossref","unstructured":"Wang, H., Dong, X., Chen, Z., & Shi, B.E. (2015). Hybrid gaze\/eeg brain computer interface for robot arm control on a pick and place task. In: 2015 37th Annual International Conference of the IEEE Engineering in Medicine and Biology Society (EMBC), pp. 1476\u20131479. IEEE","DOI":"10.1109\/EMBC.2015.7318649"},{"key":"2233_CR46","doi-asserted-by":"crossref","unstructured":"Wang, H., Wu, X., Huang, Z., & Xing, E.P. (2020). High-frequency component helps explain the generalization of convolutional neural networks. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 8684\u20138694","DOI":"10.1109\/CVPR42600.2020.00871"},{"key":"2233_CR47","doi-asserted-by":"crossref","unstructured":"Wang, K., Zhao, R., Su, H., & Ji, Q. (2019). Generalizing eye tracking with bayesian adversarial learning. In: Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition, pp. 11907\u201311916","DOI":"10.1109\/CVPR.2019.01218"},{"issue":"4","key":"2233_CR48","doi-asserted-by":"publisher","first-page":"600","DOI":"10.1109\/TIP.2003.819861","volume":"13","author":"Z Wang","year":"2004","unstructured":"Wang, Z., Bovik, A. C., Sheikh, H. R., & Simoncelli, E. P. (2004). Image quality assessment: from error visibility to structural similarity. IEEE transactions on image processing, 13(4), 600\u2013612.","journal-title":"IEEE transactions on image processing"},{"key":"2233_CR49","doi-asserted-by":"crossref","unstructured":"Wu, Z., Xiong, Y., Yu, S.X., & Lin, D. (2018). Unsupervised feature learning via non-parametric instance discrimination. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp. 3733\u20133742","DOI":"10.1109\/CVPR.2018.00393"},{"key":"2233_CR50","doi-asserted-by":"publisher","first-page":"3027","DOI":"10.1609\/aaai.v37i3.25406","volume":"37","author":"M Xu","year":"2023","unstructured":"Xu, M., Wang, H., & Lu, F. (2023). Learning a generalized gaze estimator from gaze-consistent feature. Proceedings of the AAAI Conference on Artificial Intelligence, 37, 3027\u20133035.","journal-title":"Proceedings of the AAAI Conference on Artificial Intelligence"},{"key":"2233_CR51","doi-asserted-by":"crossref","unstructured":"Yang, J., Li, C., An, W., Ma, H., Guo, Y., Rong, Y., Zhao, P., & Huang, J. (2021) Exploring robustness of unsupervised domain adaptation in semantic segmentation. arXiv preprint arXiv:2105.10843","DOI":"10.1109\/ICCV48922.2021.00906"},{"key":"2233_CR52","doi-asserted-by":"crossref","unstructured":"Zhang, X., Park, S., Beeler, T., Bradley, D., Tang, S., & Hilliges, O. (2020). Eth-xgaze: A large scale dataset for gaze estimation under extreme head pose and gaze variation. In: European Conference on Computer Vision, pp. 365\u2013381 . Springer","DOI":"10.1007\/978-3-030-58558-7_22"},{"issue":"1","key":"2233_CR53","doi-asserted-by":"publisher","first-page":"162","DOI":"10.1109\/TPAMI.2017.2778103","volume":"41","author":"X Zhang","year":"2017","unstructured":"Zhang, X., Sugano, Y., Fritz, M., & Bulling, A. (2017). Mpiigaze: Real-world dataset and deep appearance-based gaze estimation. IEEE transactions on pattern analysis and machine intelligence, 41(1), 162\u2013175.","journal-title":"IEEE transactions on pattern analysis and machine intelligence"},{"key":"2233_CR54","unstructured":"Zimmermann, R. S. (2019) Comment on\" adv-bnn: Improved adversarial defense through robust bayesian neural network\". arXiv preprint arXiv:1907.00895"}],"container-title":["International Journal of Computer Vision"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11263-024-02233-1.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11263-024-02233-1\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11263-024-02233-1.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,2,24]],"date-time":"2025-02-24T10:03:43Z","timestamp":1740391423000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11263-024-02233-1"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2024,9,30]]},"references-count":54,"journal-issue":{"issue":"3","published-print":{"date-parts":[[2025,3]]}},"alternative-id":["2233"],"URL":"https:\/\/doi.org\/10.1007\/s11263-024-02233-1","relation":{},"ISSN":["0920-5691","1573-1405"],"issn-type":[{"value":"0920-5691","type":"print"},{"value":"1573-1405","type":"electronic"}],"subject":[],"published":{"date-parts":[[2024,9,30]]},"assertion":[{"value":"7 December 2023","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"24 August 2024","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"30 September 2024","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors have no Conflict of interest to declare that are relevant to the content of this article.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interest"}}]}}