{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2025,10,29]],"date-time":"2025-10-29T13:51:00Z","timestamp":1761745860720},"reference-count":90,"publisher":"Springer Science and Business Media LLC","issue":"8","license":[{"start":{"date-parts":[[2023,8,28]],"date-time":"2023-08-28T00:00:00Z","timestamp":1693180800000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2023,8,28]],"date-time":"2023-08-28T00:00:00Z","timestamp":1693180800000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"name":"The Natural Science Foundation of Hebei Province","award":["F2019201451","F2019201451","F2019201451"],"award-info":[{"award-number":["F2019201451","F2019201451","F2019201451"]}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Neural Process Lett"],"published-print":{"date-parts":[[2023,12]]},"DOI":"10.1007\/s11063-023-11395-x","type":"journal-article","created":{"date-parts":[[2023,8,28]],"date-time":"2023-08-28T12:04:00Z","timestamp":1693224240000},"page":"11701-11719","update-policy":"http:\/\/dx.doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":1,"title":["Multi-scale Deep Feature Transfer for Automatic Video Object Segmentation"],"prefix":"10.1007","volume":"55","author":[{"given":"Zhen","family":"Yang","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Qingxuan","family":"Shi","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Yichuan","family":"Fang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2023,8,28]]},"reference":[{"key":"11395_CR1","doi-asserted-by":"crossref","unstructured":"Chen X, Li Z, Yuan Y, Yu G, Shen J, Qi D (2020) State-aware tracker for real-time video object segmentation. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 9384\u20139393","DOI":"10.1109\/CVPR42600.2020.00940"},{"key":"11395_CR2","doi-asserted-by":"crossref","unstructured":"Huang X, Xu J, Tai Y.-W, Tang C.-K (2020) Fast video object segmentation with temporal aggregation network and dynamic template matching. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 8879\u20138889","DOI":"10.1109\/CVPR42600.2020.00890"},{"key":"11395_CR3","doi-asserted-by":"crossref","unstructured":"Wang Y, Xu Z, Wang X, Shen C, Cheng B, Shen H, Xia H (2021) End-to-end video instance segmentation with transformers. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 8741\u20138750","DOI":"10.1109\/CVPR46437.2021.00863"},{"key":"11395_CR4","doi-asserted-by":"crossref","unstructured":"Liu J, Dai H.-N, Zhao G, Li B, Zhang T (2022) TMVOS: triplet matching for efficient video object segmentation. Signal Process Image Commun 107","DOI":"10.1016\/j.image.2022.116779"},{"issue":"1","key":"11395_CR5","doi-asserted-by":"publisher","first-page":"3","DOI":"10.1177\/0278364916679498","volume":"36","author":"W Maddern","year":"2017","unstructured":"Maddern W, Pascoe G, Linegar C, Newman P (2017) 1 year, 1000 km: The oxford robotcar dataset. Int J Robot Res 36(1):3\u201315","journal-title":"Int J Robot Res"},{"issue":"1","key":"11395_CR6","doi-asserted-by":"publisher","first-page":"19","DOI":"10.1109\/TIP.2013.2282897","volume":"23","author":"H Hadizadeh","year":"2013","unstructured":"Hadizadeh H, Baji\u0107 IV (2013) Saliency-aware video compression. IEEE Trans Image Process 23(1):19\u201333","journal-title":"IEEE Trans Image Process"},{"key":"11395_CR7","doi-asserted-by":"crossref","unstructured":"Chen Y, Pont-Tuset J, Montes A, Van\u00a0Gool L (2018) Blazingly fast video object segmentation with pixel-wise metric learning. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 1189\u20131198","DOI":"10.1109\/CVPR.2018.00130"},{"key":"11395_CR8","unstructured":"Zhou T, Porikli F, Crandall D.J, Van\u00a0Gool L, Wang W (2022) A survey on deep learning technique for video segmentation. In: IEEE Transactions on pattern analysis and machine intelligence. IEEE, pp 1\u201320"},{"issue":"4","key":"11395_CR9","doi-asserted-by":"publisher","first-page":"985","DOI":"10.1109\/TPAMI.2018.2819173","volume":"41","author":"W Wang","year":"2018","unstructured":"Wang W, Shen J, Porikli F, Yang R (2018) Semi-supervised video object segmentation with super-trajectories. IEEE Trans Pattern Anal Mach Intell 41(4):985\u2013998","journal-title":"IEEE Trans Pattern Anal Mach Intell"},{"key":"11395_CR10","doi-asserted-by":"crossref","unstructured":"Bhat G, Lawin F.J, Danelljan M, Robinson A, Felsberg M, Van\u00a0Gool L, Timofte R (2020) Learning what to learn for video object segmentation. In: Computer vision\u2013ECCV 2020: 16th European conference, Glasgow, UK, August 23\u201328, 2020, Proceedings, Part II 16. Springer, pp 777\u2013794","DOI":"10.1007\/978-3-030-58536-5_46"},{"key":"11395_CR11","unstructured":"Caelles S, Pont-Tuset J, Perazzi F, Montes A, Maninis K.-K, Van\u00a0Gool L (2019) The 2019 davis challenge on vos: Unsupervised multi-object segmentation. CoRR abs\/1905.00737"},{"key":"11395_CR12","doi-asserted-by":"crossref","unstructured":"Lan M, Zhang Y, Xu Q, Zhang L (2020) E3sn: efficient end-to-end siamese network for video object segmentation. In: IJCAI, pp 701\u2013707","DOI":"10.24963\/ijcai.2020\/98"},{"key":"11395_CR13","doi-asserted-by":"crossref","unstructured":"Li Y, Shen Z, Shan Y (2020) Fast video object segmentation using the global context module. In: Computer vision\u2013ECCV 2020: 16th European conference, Glasgow, UK, August 23\u201328, 2020, proceedings, Part X 16. Springer, pp 735\u2013750","DOI":"10.1007\/978-3-030-58607-2_43"},{"key":"11395_CR14","doi-asserted-by":"crossref","unstructured":"Robinson A, Lawin F.J, Danelljan M, Khan F.S, Felsberg M (2020) Learning fast and robust target models for video object segmentation, in: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 7406\u20137415","DOI":"10.1109\/CVPR42600.2020.00743"},{"key":"11395_CR15","doi-asserted-by":"crossref","unstructured":"Seong H, Hyun J, Kim E (2020) Kernelized memory network for video object segmentation. In: Computer vision\u2013ECCV 2020: 16th European conference, Glasgow, UK, August 23\u201328, 2020, proceedings, Part XXII 16. Springer, pp 629\u2013645","DOI":"10.1007\/978-3-030-58542-6_38"},{"key":"11395_CR16","doi-asserted-by":"crossref","unstructured":"Xu N, Yang L, Fan Y, Yang J, Yue D, Liang Y, Price B, Cohen S, Huang T (2018) Youtube-vos: sequence-to-sequence video object segmentation. In: Proceedings of the European conference on computer vision (ECCV), pp 585\u2013601","DOI":"10.1007\/978-3-030-01228-1_36"},{"key":"11395_CR17","doi-asserted-by":"crossref","unstructured":"Yang L, Fan Y, Xu N (2019) Video instance segmentation. In: Proceedings of the IEEE\/CVF international conference on computer vision, pp 5188\u20135197","DOI":"10.1109\/ICCV.2019.00529"},{"key":"11395_CR18","doi-asserted-by":"crossref","unstructured":"Zhang K, Wang L, Liu D, Liu B, Liu Q, Li Z (2020) Dual temporal memory network for efficient video object segmentation. In: Proceedings of the 28th ACM international conference on multimedia, pp 1515\u20131523","DOI":"10.1145\/3394171.3413942"},{"key":"11395_CR19","doi-asserted-by":"crossref","unstructured":"Zhang Y, Wu Z, Peng H, Lin S (2020) A transductive approach for video object segmentation. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 6949\u20136958","DOI":"10.1109\/CVPR42600.2020.00698"},{"key":"11395_CR20","unstructured":"Mahadevan S, Athar A, O\u0161ep A, Hennen S, Leal-Taix\u00e9 L, Leibe B (2020) Making a case for 3d convolutions for object segmentation in videos. CoRR abs\/2008.11516"},{"key":"11395_CR21","doi-asserted-by":"crossref","unstructured":"Lu X, Wang W, Ma C, Shen J, Shao L, Porikli F (2019) See more, know more: unsupervised video object segmentation with co-attention siamese networks. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 3623\u20133632","DOI":"10.1109\/CVPR.2019.00374"},{"key":"11395_CR22","doi-asserted-by":"crossref","unstructured":"Tokmakov P, Alahari K, Schmid C (2017) Learning video object segmentation with visual memory. In: Proceedings of the IEEE international conference on computer vision, pp 4481\u20134490","DOI":"10.1109\/ICCV.2017.480"},{"key":"11395_CR23","doi-asserted-by":"crossref","unstructured":"Tokmakov P, Schmid C, Alahari K (2019) Learning to segment moving objects. In: International journal of computer vision. Springer, pp 282\u2013301","DOI":"10.1007\/s11263-018-1122-2"},{"key":"11395_CR24","doi-asserted-by":"crossref","unstructured":"Yang Z, Wang Q, Bertinetto L, Hu W, Bai S, Torr P.H (2019) Anchor diffusion for unsupervised video object segmentation. In: Proceedings of the IEEE\/CVF international conference on computer vision, pp 931\u2013940","DOI":"10.1109\/ICCV.2019.00102"},{"key":"11395_CR25","doi-asserted-by":"crossref","unstructured":"Li G, Xie Y, Lin L, Yu Y (2017) Instance-level salient object segmentation. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 2386\u20132395","DOI":"10.1109\/CVPR.2017.34"},{"key":"11395_CR26","doi-asserted-by":"crossref","unstructured":"Hou Q, Cheng M.-M, Hu X, Borji A, Tu Z, Torr P.H (2017) Deeply supervised salient object detection with short connections. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 3203\u20133212","DOI":"10.1109\/CVPR.2017.563"},{"key":"11395_CR27","doi-asserted-by":"crossref","unstructured":"Li G, Yu Y (2016) Deep contrast learning for salient object detection. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 478\u2013487","DOI":"10.1109\/CVPR.2016.58"},{"key":"11395_CR28","doi-asserted-by":"crossref","unstructured":"Wang W, Shen J (2017) Deep visual attention prediction. In: IEEE Transactions on image processing. IEEE, pp 2368\u20132378","DOI":"10.1109\/TIP.2017.2787612"},{"key":"11395_CR29","doi-asserted-by":"crossref","unstructured":"Li H, Chen G, Li G, Yu Y (2019) Motion guided attention for video salient object detection. In: Proceedings of the IEEE\/CVF international conference on computer vision, pp 7274\u20137283","DOI":"10.1109\/ICCV.2019.00737"},{"key":"11395_CR30","doi-asserted-by":"crossref","unstructured":"Tokmakov P, Alahari K, Schmid C (2017) Learning motion patterns in videos. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 3386\u20133394","DOI":"10.1109\/CVPR.2017.64"},{"key":"11395_CR31","doi-asserted-by":"crossref","unstructured":"Perazzi F, Khoreva A, Benenson R, Schiele B, Sorkine-Hornung A (2017) Learning video object segmentation from static images. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 2663\u20132672","DOI":"10.1109\/CVPR.2017.372"},{"key":"11395_CR32","doi-asserted-by":"crossref","unstructured":"Dutt\u00a0Jain S, Xiong B, Grauman K (2017) Fusionseg: learning to combine motion and appearance for fully automatic segmentation of generic objects in videos. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 3664\u20133673","DOI":"10.1109\/CVPR.2017.228"},{"key":"11395_CR33","doi-asserted-by":"crossref","unstructured":"Cheng J, Tsai Y.-H, Wang S, Yang M.-H (2017) Segflow: joint learning for video object segmentation and optical flow. In: Proceedings of the IEEE international conference on computer vision, pp 686\u2013695","DOI":"10.1109\/ICCV.2017.81"},{"key":"11395_CR34","doi-asserted-by":"crossref","unstructured":"Li S, Seybold B, Vorobyov A, Lei X, Kuo C.-C.J (2018) Unsupervised video object segmentation with motion-based bilateral networks. In: Proceedings of the European conference on computer vision (ECCV), pp 207\u2013223","DOI":"10.1007\/978-3-030-01219-9_13"},{"key":"11395_CR35","doi-asserted-by":"crossref","unstructured":"Zhou T, Li J, Wang S, Tao R, Shen J (2020) Matnet: motion-attentive transition network for zero-shot video object segmentation. In: IEEE Transactions on image processing, pp 8326\u20138338","DOI":"10.1109\/TIP.2020.3013162"},{"key":"11395_CR36","doi-asserted-by":"crossref","unstructured":"Ji G.-P, Fu K, Wu Z, Fan D.-P, Shen J, Shao L (2021) Full-duplex strategy for video object segmentation. In: Proceedings of the IEEE\/CVF international conference on computer vision, pp 4922\u20134933","DOI":"10.1109\/ICCV48922.2021.00488"},{"key":"11395_CR37","doi-asserted-by":"crossref","unstructured":"Perazzi F, Pont-Tuset J, McWilliams B, Van\u00a0Gool L, Gross M, Sorkine-Hornung A (2016) A benchmark dataset and evaluation methodology for video object segmentation. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 724\u2013732","DOI":"10.1109\/CVPR.2016.85"},{"key":"11395_CR38","doi-asserted-by":"crossref","unstructured":"Ochs P, Malik J, Brox T (2013) Segmentation of moving objects by long term video analysis. In: IEEE Transaction on pattern analysis and machine intelligence, pp 1187\u20131200","DOI":"10.1109\/TPAMI.2013.242"},{"key":"11395_CR39","doi-asserted-by":"crossref","unstructured":"Tsai Y.-H, Yang M.-H, Black M.J (2016) Video segmentation via object flow. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 3899\u20133908","DOI":"10.1109\/CVPR.2016.423"},{"key":"11395_CR40","doi-asserted-by":"crossref","unstructured":"Xu Y.-S, Fu T.-J, Yang H.-K, Lee C.-Y (2018) Dynamic video segmentation network. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 6556\u20136565","DOI":"10.1109\/CVPR.2018.00686"},{"key":"11395_CR41","doi-asserted-by":"crossref","unstructured":"Wang J, Chen D, Wu Z, Luo C, Tang C, Dai X, Zhao Y, Xie Y, Yuan L, Jiang Y.-G (2023) Look before you match: instance understanding matters in video object segmentation. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 2268\u20132278","DOI":"10.1109\/CVPR52729.2023.00225"},{"key":"11395_CR42","doi-asserted-by":"crossref","unstructured":"Cheng H.K, Schwing A.G (2022) Xmem: long-term video object segmentation with an Atkinson\u2013Shiffrin memory model. In: European conference on computer vision. Springer, pp 640\u2013658","DOI":"10.1007\/978-3-031-19815-1_37"},{"key":"11395_CR43","doi-asserted-by":"crossref","unstructured":"Hu Y.-T, Huang J.-B, Schwing A.G (2018) Unsupervised video object segmentation using motion saliency-guided spatio-temporal propagation. In: Proceedings of the European conference on computer vision (ECCV), pp 786\u2013802","DOI":"10.1007\/978-3-030-01246-5_48"},{"key":"11395_CR44","doi-asserted-by":"publisher","first-page":"3137","DOI":"10.1109\/TIP.2015.2438550","volume":"24","author":"W Wang","year":"2015","unstructured":"Wang W, Shen J, Li X, Porikli F (2015) Robust video object cosegmentation. IEEE Trans Image Process 24:3137\u20133148","journal-title":"IEEE Trans Image Process"},{"key":"11395_CR45","doi-asserted-by":"crossref","unstructured":"Wang W, Shen J, Porikli F (2015) Saliency-aware geodesic video object segmentation. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 3395\u20133402","DOI":"10.1109\/CVPR.2015.7298961"},{"key":"11395_CR46","doi-asserted-by":"crossref","unstructured":"Faktor A, Irani M (2014) Video segmentation by non-local consensus voting. In: BMVC, p 8","DOI":"10.5244\/C.28.21"},{"key":"11395_CR47","doi-asserted-by":"crossref","unstructured":"Lee Y.J, Kim J, Grauman K (2011) Key-segments for video object segmentation. In: 2011 International conference on computer vision. IEEE, pp 1995\u20132002","DOI":"10.1109\/ICCV.2011.6126471"},{"key":"11395_CR48","doi-asserted-by":"crossref","unstructured":"Li F, Kim T, Humayun A, Tsai D, Rehg J.M (2013) Video segmentation by tracking many figure-ground segments. In: Proceedings of the IEEE international conference on computer vision, pp 2192\u20132199","DOI":"10.1109\/ICCV.2013.273"},{"key":"11395_CR49","doi-asserted-by":"crossref","unstructured":"Robinson A, Lawin F.J, Danelljan M, Khan F.S, Felsberg M (2020) Learning fast and robust target models for video object segmentation. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 7406\u20137415","DOI":"10.1109\/CVPR42600.2020.00743"},{"key":"11395_CR50","unstructured":"Ballas N, Yao L, Pal C, Courville A (2016) Delving deeper into convolutional networks for learning video representations"},{"key":"11395_CR51","doi-asserted-by":"crossref","unstructured":"Song H, Wang W, Zhao S, Shen J, Lam K.-M (2018) Pyramid dilated deeper ConvLSTM for video salient object detection. In: Proceedings of the European conference on computer vision (ECCV), pp 715\u2013731","DOI":"10.1007\/978-3-030-01252-6_44"},{"key":"11395_CR52","doi-asserted-by":"crossref","unstructured":"Wang W, Song H, Zhao S, Shen J, Zhao S, Hoi S.\u00a0C, Ling H (2019) Learning unsupervised video object segmentation through visual attention. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 3064\u20133074","DOI":"10.1109\/CVPR.2019.00318"},{"key":"11395_CR53","first-page":"2191","volume":"30","author":"M Xu","year":"2019","unstructured":"Xu M, Liu B, Fu P, Li J, Hu YH, Feng S (2019) Video salient object detection via robust seeds extraction and multi-graphs manifold propagation. IEEE Trans Circuits Syst Video Technol 30:2191\u20132206","journal-title":"IEEE Trans Circuits Syst Video Technol"},{"key":"11395_CR54","doi-asserted-by":"crossref","unstructured":"Zheng J, Luo W, Piao Z (2019) Cascaded ConvLSTMs using semantically-coherent data synthesis for video object segmentation. In: IEEE access, pp 132120\u2013132129","DOI":"10.1109\/ACCESS.2019.2940768"},{"key":"11395_CR55","unstructured":"Simonyan K, Zisserman A (2014) Two-stream convolutional networks for action recognition in videos. Adv Neural Inf Process Syst 27"},{"key":"11395_CR56","doi-asserted-by":"crossref","unstructured":"Wang W, Lu X, Shen J, Crandall D.J, Shao L (2019) Zero-shot video object segmentation via attentive graph neural networks. In: Proceedings of the IEEE\/CVF international conference on computer vision, pp 9236\u20139245","DOI":"10.1109\/ICCV.2019.00933"},{"key":"11395_CR57","doi-asserted-by":"crossref","unstructured":"Galasso F, Cipolla R, Schiele B (2013) Video segmentation with superpixels. In: Computer vision\u2013ACCV 2012: 11th Asian conference on computer vision, Daejeon, Korea, November 5\u20139, 2012, Revised Selected Papers, Part I 11. Springer, pp 760\u2013774","DOI":"10.1007\/978-3-642-37331-2_57"},{"key":"11395_CR58","doi-asserted-by":"crossref","unstructured":"Grundmann M, Kwatra V, Han M, Essa I (2010) Efficient hierarchical graph-based video segmentation. In: 2010 IEEE Computer society conference on computer vision and pattern recognition. IEEE, pp 2141\u20132148","DOI":"10.1109\/CVPR.2010.5539893"},{"key":"11395_CR59","doi-asserted-by":"crossref","unstructured":"Xu C, Xiong C, Corso J.J (2012) Streaming hierarchical video segmentation. In: Computer vision\u2013ECCV 2012: 12th European conference on computer vision, Florence, Italy, October 7\u201313, 2012, Proceedings, Part VI 12. Springer, pp 626\u2013639","DOI":"10.1007\/978-3-642-33783-3_45"},{"key":"11395_CR60","doi-asserted-by":"crossref","unstructured":"Li X, Loy C.C (2018) Video object segmentation with joint re-identification and attention-aware mask propagation. In: Proceedings of the European conference on computer vision (ECCV), pp 90\u2013105","DOI":"10.1007\/978-3-030-01219-9_6"},{"key":"11395_CR61","doi-asserted-by":"crossref","unstructured":"He K, Zhang X, Ren S, Sun J (2016) Deep residual learning for image recognition. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 770\u2013778","DOI":"10.1109\/CVPR.2016.90"},{"key":"11395_CR62","doi-asserted-by":"crossref","unstructured":"Long J, Shelhamer E, Darrell T (2015) Fully convolutional networks for semantic segmentation. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 3431\u20133440","DOI":"10.1109\/CVPR.2015.7298965"},{"key":"11395_CR63","doi-asserted-by":"crossref","unstructured":"Zhao H, Shi J, Qi X, Wang X, Jia J (2017) Pyramid scene parsing network. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 2881\u20132890","DOI":"10.1109\/CVPR.2017.660"},{"key":"11395_CR64","unstructured":"Paszke A, Gross S, Massa F, Lerer A, Bradbury J, Chanan G, Killeen T, Lin Z, Gimelshein N, Antiga L et al (2019) Pytorch: An imperative style, high-performance deep learning library. Adv Neural Inf Process Syst 32"},{"key":"11395_CR65","doi-asserted-by":"crossref","unstructured":"Wang L, Lu H, Wang Y, Feng M, Wang D, Yin B, Ruan X (2017) Learning to detect salient objects with image-level supervision. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 136\u2013145","DOI":"10.1109\/CVPR.2017.404"},{"key":"11395_CR66","unstructured":"Kr\u00e4henb\u00fchl P, Koltun V (2011) Efficient inference in fully connected CRFS with gaussian edge potentials. Adv Neural Inf Process Syst 24"},{"key":"11395_CR67","doi-asserted-by":"crossref","unstructured":"Papazoglou A, Ferrari V (2013) Fast object segmentation in unconstrained video. In: Proceedings of the IEEE international conference on computer vision, pp 1777\u20131784","DOI":"10.1109\/ICCV.2013.223"},{"key":"11395_CR68","doi-asserted-by":"crossref","unstructured":"Lao D, Sundaramoorthi G (2018) Extending layered models to 3d motion. In: Proceedings of the European conference on computer vision (ECCV), pp 435\u2013451","DOI":"10.1007\/978-3-030-01249-6_27"},{"key":"11395_CR69","doi-asserted-by":"crossref","unstructured":"Tokmakov P, Alahari K, Schmid C (2017) Learning video object segmentation with visual memory. In: Proceedings of the IEEE international conference on computer vision, pp 4481\u20134490","DOI":"10.1109\/ICCV.2017.480"},{"key":"11395_CR70","doi-asserted-by":"crossref","unstructured":"Koh Y.\u00a0J, Kim C.-S (2017) Primary object segmentation in videos based on region augmentation and reduction. In: 2017 IEEE Conference on computer vision and pattern recognition (CVPR). IEEE, pp 7417\u20137425","DOI":"10.1109\/CVPR.2017.784"},{"key":"11395_CR71","doi-asserted-by":"crossref","unstructured":"Siam M, Jiang C, Lu S, Petrich L, Gamal M, Elhoseiny M, Jagersand M (2019) Video object segmentation using teacher-student adaptation in a human robot interaction (HRI) setting. In: 2019 International conference on robotics and automation (ICRA). IEEE, pp 50\u201356","DOI":"10.1109\/ICRA.2019.8794254"},{"key":"11395_CR72","unstructured":"Akhter I, Ali M, Faisal M, Hartley R (2020) Epo-net: exploiting geometric constraints on dense trajectories for motion saliency. In: Proceedings of the IEEE\/CVF winter conference on applications of computer vision, pp 1884\u20131893"},{"key":"11395_CR73","doi-asserted-by":"crossref","unstructured":"Chen Y.-W, Jin X, Shen X, Yang M.-H (2022) Video salient object detection via contrastive features and attention modules. In: Proceedings of the IEEE\/CVF winter conference on applications of computer vision, pp 1320\u20131329","DOI":"10.1109\/WACV51458.2022.00061"},{"key":"11395_CR74","doi-asserted-by":"crossref","unstructured":"Lee M, Cho S, Lee S, Park C, Lee S (2023) Unsupervised video object segmentation via prototype memory network. In: Proceedings of the IEEE\/CVF winter conference on applications of computer vision, pp 5924\u20135934","DOI":"10.1109\/WACV56688.2023.00587"},{"key":"11395_CR75","doi-asserted-by":"crossref","unstructured":"Fan D.-P, Cheng M.-M, Liu Y, Li T, Borji A (2017) Structure-measure: a new way to evaluate foreground maps. In: Proceedings of the IEEE international conference on computer vision, pp 4548\u20134557","DOI":"10.1109\/ICCV.2017.487"},{"issue":"9","key":"11395_CR76","doi-asserted-by":"publisher","first-page":"1475","DOI":"10.1360\/SSI-2020-0370","volume":"51","author":"D-P Fan","year":"2021","unstructured":"Fan D-P, Ji G-P, Qin X-B, Cheng M-M (2021) Cognitive vision inspired object segmentation metric and loss function. Sci Sin Inf 51(9):1475","journal-title":"Sci Sin Inf"},{"key":"11395_CR77","doi-asserted-by":"crossref","unstructured":"Achanta R, HemamiS, Estrada F, Susstrunk S (2009) Frequency-tuned salient region detection. In: 2009 IEEE Conference on computer vision and pattern recognition. IEEE, pp 1597\u20131604","DOI":"10.1109\/CVPR.2009.5206596"},{"key":"11395_CR78","doi-asserted-by":"crossref","unstructured":"Perazzi F, Kr\u00e4henb\u00fchl P, Pritch Y, Hornung A (2012) Saliency filters: contrast based filtering for salient region detection. In: 2012 IEEE Conference on computer vision and pattern recognition. IEEE, pp 733\u2013740","DOI":"10.1109\/CVPR.2012.6247743"},{"key":"11395_CR79","doi-asserted-by":"crossref","unstructured":"Ding M, Wang Z, Zhou B, Shi J, Lu Z, Luo P (2020) Every frame counts: joint learning of video segmentation and optical flow. In: Proceedings of the AAAI conference on artificial intelligence, pp 10713\u201310720","DOI":"10.1609\/aaai.v34i07.6699"},{"key":"11395_CR80","doi-asserted-by":"publisher","first-page":"2790","DOI":"10.1109\/TMM.2019.2914889","volume":"21","author":"M Xu","year":"2019","unstructured":"Xu M, Liu B, Fu P, Li J, Hu YH (2019) Video saliency detection via graph clustering with motion energy and spatiotemporal objectness. IEEE Trans Multimed 21:2790\u20132805","journal-title":"IEEE Trans Multimed"},{"key":"11395_CR81","doi-asserted-by":"publisher","first-page":"1973","DOI":"10.1109\/TCSVT.2018.2859773","volume":"29","author":"Y Tang","year":"2018","unstructured":"Tang Y, Zou W, Jin Z, Chen Y, Hua Y, Li X (2018) Weakly supervised salient object detection with spatiotemporal cascade neural networks. IEEE Trans Circuits Syst Video Technol 29:1973\u20131984","journal-title":"IEEE Trans Circuits Syst Video Technol"},{"key":"11395_CR82","doi-asserted-by":"crossref","unstructured":"Li Y, Li S, Chen C, Hao A, Qin H (2019) Accurate and robust video saliency detection via self-paced diffusion. In: IEEE Transactions on multimedia, pp 1153\u20131167","DOI":"10.1109\/TMM.2019.2940851"},{"key":"11395_CR83","doi-asserted-by":"crossref","unstructured":"Wang W, Shen J, Shao L (2017) Video salient object detection via fully convolutional networks. In: IEEE Transactions on image processing, pp 38\u201349","DOI":"10.1109\/TIP.2017.2754941"},{"key":"11395_CR84","doi-asserted-by":"crossref","unstructured":"Chen Y, Zou W, Tang Y, Li X, Xu C, Komodakis N (2018) SCOM: spatiotemporal constrained optimization for salient object detection. In: IEEE Transactions on image processing, pp 3345\u20133357","DOI":"10.1109\/TIP.2018.2813165"},{"key":"11395_CR85","doi-asserted-by":"crossref","unstructured":"Li G, Xie Y, Wei T, Wang K, Lin L (2018) Flow guided recurrent neural encoder for video salient object detection. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 3243\u20133252","DOI":"10.1109\/CVPR.2018.00342"},{"key":"11395_CR86","doi-asserted-by":"crossref","unstructured":"Chen C, Wang G, Peng C, Zhang X, Qin H (2019) Improved robust video saliency detection based on long-term spatial-temporal information. In: IEEE Transactions on image processing, pp 1090\u20131100","DOI":"10.1109\/TIP.2019.2934350"},{"key":"11395_CR87","doi-asserted-by":"crossref","unstructured":"Yan P, Li G, Xie Y, Li Z, Wang C, Chen T , Lin L (2019) Semi-supervised video salient object detection using pseudo-labels. In: Proceedings of the IEEE\/CVF international conference on computer vision, pp 7284\u20137293","DOI":"10.1109\/ICCV.2019.00738"},{"key":"11395_CR88","doi-asserted-by":"crossref","unstructured":"Fan D.-P, Wang W, Cheng M.-M, Shen J (2019) Shifting more attention to video salient object detection. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition, pp 8554\u20138564","DOI":"10.1109\/CVPR.2019.00875"},{"key":"11395_CR89","doi-asserted-by":"crossref","unstructured":"Gu Y, Wang L, Wang Z, Liu Y, Cheng M.-M, Lu S.-P (2020) Pyramid constrained self-attention network for fast video salient object detection. In: Proceedings of the AAAI conference on artificial intelligence, pp 10869\u201310876","DOI":"10.1609\/aaai.v34i07.6718"},{"key":"11395_CR90","unstructured":"Shi X, Chen Z, Wang H, Yeung D.-Y, Wong W.-K, Woo W.-C (2015) Convolutional LSTM network: a machine learning approach for precipitation nowcasting. Adv Neural Inf Process Syst 28"}],"container-title":["Neural Processing Letters"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11063-023-11395-x.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11063-023-11395-x\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11063-023-11395-x.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2023,11,22]],"date-time":"2023-11-22T05:20:08Z","timestamp":1700630408000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11063-023-11395-x"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2023,8,28]]},"references-count":90,"journal-issue":{"issue":"8","published-print":{"date-parts":[[2023,12]]}},"alternative-id":["11395"],"URL":"https:\/\/doi.org\/10.1007\/s11063-023-11395-x","relation":{},"ISSN":["1370-4621","1573-773X"],"issn-type":[{"value":"1370-4621","type":"print"},{"value":"1573-773X","type":"electronic"}],"subject":[],"published":{"date-parts":[[2023,8,28]]},"assertion":[{"value":"15 August 2023","order":1,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"28 August 2023","order":2,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflicts of interest"}}]}}