{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2025,7,4]],"date-time":"2025-07-04T05:39:29Z","timestamp":1751607569289,"version":"3.37.3"},"reference-count":43,"publisher":"Springer Science and Business Media LLC","issue":"12","license":[{"start":{"date-parts":[[2023,9,27]],"date-time":"2023-09-27T00:00:00Z","timestamp":1695772800000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2023,9,27]],"date-time":"2023-09-27T00:00:00Z","timestamp":1695772800000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"funder":[{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"publisher","award":["No.U21A20518","No.61976086","No.62106071"],"award-info":[{"award-number":["No.U21A20518","No.61976086","No.62106071"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Multimed Tools Appl"],"DOI":"10.1007\/s11042-023-16913-6","type":"journal-article","created":{"date-parts":[[2023,9,27]],"date-time":"2023-09-27T08:02:34Z","timestamp":1695801754000},"page":"36287-36305","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":1,"title":["Multiscale deep feature selection fusion network for referring image segmentation"],"prefix":"10.1007","volume":"83","author":[{"given":"Xianwen","family":"Dai","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Jiacheng","family":"Lin","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Ke","family":"Nai","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Qingpeng","family":"Li","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0001-9720-5915","authenticated-orcid":false,"given":"Zhiyong","family":"Li","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2023,9,27]]},"reference":[{"key":"16913_CR1","doi-asserted-by":"publisher","first-page":"132","DOI":"10.1016\/j.neunet.2020.09.001","volume":"133","author":"J Lin","year":"2021","unstructured":"Lin J, Li Y, Yang G (2021) Fpgan: Face de-identification method with generative adversarial networks for social robots. Neural Netw 133:132\u2013147","journal-title":"Neural Netw"},{"key":"16913_CR2","doi-asserted-by":"crossref","unstructured":"Zhao H, Shi J, Qi X, Wang X, Jia J (2017) Pyramid scene parsing network. In: Proceedings of the IEEE conference on computer vision and pattern recognition (CVPR)","DOI":"10.1109\/CVPR.2017.660"},{"key":"16913_CR3","unstructured":"Zhang H, Wu C, Zhang Z, Zhu Y, Zhang Z, Lin H, Sun Y, He T, Mueller J, Manmatha R, Li M, Smola AJ (2020) Resnest: Split-attention networks. CoRR arXiv:2004.08955"},{"key":"16913_CR4","doi-asserted-by":"publisher","first-page":"18689","DOI":"10.1007\/s11042-018-5653-x","volume":"77","author":"DM Vo","year":"2018","unstructured":"Vo DM, Lee S-W (2018) Semantic image segmentation using fully convolutional neural networks with multi-scale images and multi-scale dilated convolutions. Multimed Tools Appl 77:18689\u201318707","journal-title":"Multimed Tools Appl"},{"key":"16913_CR5","doi-asserted-by":"crossref","unstructured":"Long J, Shelhamer E, Darrell T (2015) Fully convolutional networks for semantic segmentation. In: Proceedings of the IEEE conference on computer vision and pattern recognition (CVPR)","DOI":"10.1109\/CVPR.2015.7298965"},{"issue":"4","key":"16913_CR6","doi-asserted-by":"publisher","first-page":"834","DOI":"10.1109\/TPAMI.2017.2699184","volume":"40","author":"L-C Chen","year":"2018","unstructured":"Chen L-C, Papandreou G, Kokkinos I, Murphy K, Yuille AL (2018) Deeplab: Semantic image segmentation with deep convolutional nets, atrous convolution, and fully connected crfs. IEEE Trans Pattern Anal Mach Intell 40(4):834\u2013848. https:\/\/doi.org\/10.1109\/TPAMI.2017.2699184","journal-title":"IEEE Trans Pattern Anal Mach Intell"},{"key":"16913_CR7","doi-asserted-by":"crossref","unstructured":"Ding H, Jiang X, Shuai B, Liu AQ, Wang G (2018) Context contrasted feature and gated multi-scale aggregation for scene segmentation. In: Proceedings of the IEEE conference on computer vision and pattern recognition (CVPR)","DOI":"10.1109\/CVPR.2018.00254"},{"issue":"4","key":"16913_CR8","doi-asserted-by":"publisher","first-page":"1731","DOI":"10.1109\/TCYB.2020.2969046","volume":"51","author":"J Yu","year":"2021","unstructured":"Yu J, Yao J, Zhang J, Yu Z, Tao D (2021) Sprnet: Single-pixel reconstruction for one-stage instance segmentation. IEEE Trans Cybernet 51(4):1731\u20131742. https:\/\/doi.org\/10.1109\/TCYB.2020.2969046","journal-title":"IEEE Trans Cybernet"},{"key":"16913_CR9","doi-asserted-by":"publisher","unstructured":"Yin C, Tang J, Yuan T, Xu Z, Wang Y (2021) Bridging the gap between semantic segmentation and instance segmentation. IEEE Trans Multimed 1\u20131. https:\/\/doi.org\/10.1109\/TMM.2021.3114541","DOI":"10.1109\/TMM.2021.3114541"},{"issue":"2","key":"16913_CR10","doi-asserted-by":"publisher","first-page":"457","DOI":"10.1109\/TMM.2018.2859746","volume":"21","author":"K Fu","year":"2019","unstructured":"Fu K, Zhao Q (2019) Gu IY-H: Refinet: A deep segmentation assisted refinement network for salient object detection. IEEE Trans Multimed 21(2):457\u2013469. https:\/\/doi.org\/10.1109\/TMM.2018.2859746","journal-title":"IEEE Trans Multimed"},{"key":"16913_CR11","doi-asserted-by":"publisher","first-page":"114428","DOI":"10.1016\/j.eswa.2020.114428","volume":"168","author":"M Moradi","year":"2021","unstructured":"Moradi M, Bayat F (2021) A salient object segmentation framework using diffusion-based affinity learning. Expert Syst Appl 168:114428. https:\/\/doi.org\/10.1016\/j.eswa.2020.114428","journal-title":"Expert Syst Appl"},{"key":"16913_CR12","doi-asserted-by":"crossref","unstructured":"Margffoy-Tuay E, P\u00e9rez J.C, Botero E, Arbel\u00e1ez P (2018) Dynamic multimodal instance segmentation guided by natural language queries. In: Proceedings of the European conference on computer vision (ECCV). pp 630\u2013645","DOI":"10.1007\/978-3-030-01252-6_39"},{"key":"16913_CR13","doi-asserted-by":"crossref","unstructured":"Shi H, Li H, Meng F, Wu Q (2018) Key-word-aware network for referring expression image segmentation. In: Proceedings of the European conference on computer vision (ECCV). pp 38\u201354","DOI":"10.1007\/978-3-030-01231-1_3"},{"key":"16913_CR14","unstructured":"Chen J, Lin J, Xiao Z, Fu H, Nai K, Yang K, Li Z (2023) EPCFormer: expression prompt collaboration transformer for universal referring video object segmentation"},{"key":"16913_CR15","doi-asserted-by":"crossref","unstructured":"Hu R, Rohrbach M, Darrell T (2016) Segmentation from natural language expressions. In: European conference on computer vision (ECCV). Springer, pp 108\u2013124","DOI":"10.1007\/978-3-319-46448-0_7"},{"key":"16913_CR16","doi-asserted-by":"crossref","unstructured":"Li R, Li K, Kuo Y-C, Shu M, Qi X, Shen X, Jia J (2018) Referring image segmentation via recurrent refinement networks. In: Proceedings of the IEEE conference on computer vision and pattern recognition (CVPR)","DOI":"10.1109\/CVPR.2018.00602"},{"key":"16913_CR17","doi-asserted-by":"crossref","unstructured":"Liu C, Lin Z, Shen X, Yang J, Lu X, Yuille A (2017) Recurrent multimodal interaction for referring image segmentation. In: Proceedings of the IEEE international conference on computer vision (ICCV)","DOI":"10.1109\/ICCV.2017.143"},{"key":"16913_CR18","doi-asserted-by":"publisher","first-page":"120960","DOI":"10.1016\/j.eswa.2023.120960","volume":"233","author":"J Lin","year":"2023","unstructured":"Lin J, Dai X, Nai K, Yuan J, Li Z, Zhang X, Li S (2023) Brppnet: Balanced privacy protection network for referring personal image privacy protection. Expert Syst Appl 233:120960","journal-title":"Expert Syst Appl"},{"key":"16913_CR19","doi-asserted-by":"crossref","unstructured":"Chen D-J, Jia S, Lo Y-C, Chen H-T, Liu T-L (2019) See-through-text grouping for referring image segmentation. In: Proceedings of the IEEE\/CVF international conference on computer vision (ICCV)","DOI":"10.1109\/ICCV.2019.00755"},{"key":"16913_CR20","doi-asserted-by":"publisher","unstructured":"Feng G, Hu Z, Zhang L, Sun J, Lu H (2021) Bidirectional relationship inferring network for referring image localization and segmentation. IEEE Trans Neural Netw Learn Sys 1\u201313. https:\/\/doi.org\/10.1109\/TNNLS.2021.3106153","DOI":"10.1109\/TNNLS.2021.3106153"},{"key":"16913_CR21","doi-asserted-by":"crossref","unstructured":"Huang S, Hui T, Liu S, Li G, Wei Y, Han J, Liu L, Li B (2020) Referring image segmentation via cross-modal progressive comprehension. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (CVPR)","DOI":"10.1109\/CVPR42600.2020.01050"},{"key":"16913_CR22","doi-asserted-by":"crossref","unstructured":"Hui T, Liu S, Huang S, Li G, Yu S, Zhang F, Han J (2020) Linguistic structure guided context modeling for referring image segmentation. In: European conference on computer vision (ECCV). Springer, pp 59\u201375","DOI":"10.1007\/978-3-030-58607-2_4"},{"key":"16913_CR23","doi-asserted-by":"crossref","unstructured":"Ye L, Rochan M, Liu Z, Wang Y (2019) Cross-modal self-attention network for referring image segmentation. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (CVPR)","DOI":"10.1109\/CVPR.2019.01075"},{"key":"16913_CR24","first-page":"109","volume":"24","author":"P Kr\u00e4henb\u00fchl","year":"2011","unstructured":"Kr\u00e4henb\u00fchl P, Koltun V (2011) Efficient inference in fully connected crfs with gaussian edge potentials. Adv Neural Inf Process 24:109\u2013117","journal-title":"Adv Neural Inf Process"},{"key":"16913_CR25","unstructured":"Chen L, Papandreou G, Schroff F, Adam H (2017) Rethinking atrous convolution for semantic image segmentation. CoRR arXiv:1706.05587"},{"key":"16913_CR26","doi-asserted-by":"crossref","unstructured":"Ronneberger O, Fischer P, Brox T (2015) U-net: Convolutional networks for biomedical image segmentation. In: International conference on medical image computing and computer-assisted intervention. Springer, pp 234\u2013241","DOI":"10.1007\/978-3-319-24574-4_28"},{"issue":"12","key":"16913_CR27","doi-asserted-by":"publisher","first-page":"3224","DOI":"10.1109\/TMM.2020.2971171","volume":"22","author":"L Ye","year":"2020","unstructured":"Ye L, Liu Z, Wang Y (2020) Dual convolutional lstm network for referring image segmentation. IEEE Trans Multimed 22(12):3224\u20133235. https:\/\/doi.org\/10.1109\/TMM.2020.2971171","journal-title":"IEEE Trans Multimed"},{"key":"16913_CR28","doi-asserted-by":"crossref","unstructured":"Luo G, Zhou Y, Sun X, Cao L, Wu C, Deng C, Ji R (2020) Multi-task collaborative network for joint referring expression comprehension and segmentation. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (CVPR)","DOI":"10.1109\/CVPR42600.2020.01005"},{"key":"16913_CR29","doi-asserted-by":"crossref","unstructured":"Feng G, Hu Z, Zhang L, Lu H (2021) Encoder fusion network with co-attention embedding for referring image segmentation. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (CVPR). pp 15506\u201315515","DOI":"10.1109\/CVPR46437.2021.01525"},{"key":"16913_CR30","doi-asserted-by":"crossref","unstructured":"Ding H, Liu C, Wang S , Jiang X (2021)Vision-language transformer and query generation for referring segmentation. In: Proceedings of the IEEE\/CVF international conference on computer vision (ICCV). pp 16321\u201316330","DOI":"10.1109\/ICCV48922.2021.01601"},{"key":"16913_CR31","doi-asserted-by":"publisher","unstructured":"Liu C, Jiang X, Ding H (2022) Instance-specific feature propagation for referring segmentation. IEEE Trans Multimed 1\u20131. https:\/\/doi.org\/10.1109\/TMM.2022.3163578","DOI":"10.1109\/TMM.2022.3163578"},{"key":"16913_CR32","doi-asserted-by":"publisher","first-page":"99","DOI":"10.1016\/j.neucom.2021.09.066","volume":"467","author":"Q Li","year":"2022","unstructured":"Li Q, Zhang Y, Sun S, Wu J, Zhao X, Tan M (2022) Cross-modality synergy network for referring expression comprehension and segmentation. Neurocomputing 467:99\u2013114. https:\/\/doi.org\/10.1016\/j.neucom.2021.09.066","journal-title":"Neurocomputing"},{"key":"16913_CR33","doi-asserted-by":"crossref","unstructured":"Kim N, Kim D, Lan C, Zeng W, Kwak S (2022) Restr: Convolution-free referring image segmentation using transformers. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (CVPR). pp 18145\u201318154","DOI":"10.1109\/CVPR52688.2022.01761"},{"key":"16913_CR34","unstructured":"Redmon J, Farhadi A (2018) Yolov3: An incremental improvement. CoRR arXiv:1804.02767"},{"key":"16913_CR35","doi-asserted-by":"crossref","unstructured":"Pennington J, Socher R, Manning CD (2014) Glove: Global vectors for word representation. In: Proceedings of the 2014 conference on empirical methods in natural language processing (EMNLP). pp 1532\u20131543","DOI":"10.3115\/v1\/D14-1162"},{"key":"16913_CR36","unstructured":"Chung J, G\u00fcl\u00e7ehre \u00c7, Cho K, Bengio Y (2014) Empirical evaluation of gated recurrent neural networks on sequence modeling. CoRR arXiv:1412.3555"},{"key":"16913_CR37","doi-asserted-by":"crossref","unstructured":"Jing Y, Kong T, Wang W, Wang L, Li L, Tan T (2021) Locate then segment: A strong pipeline for referring image segmentation. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (CVPR). pp 9858\u20139867","DOI":"10.1109\/CVPR46437.2021.00973"},{"key":"16913_CR38","doi-asserted-by":"crossref","unstructured":"Yu L, Poirson P, Yang S, Berg AC, Berg TL (2016) Modeling context in referring expressions. In: European conference on computer vision (ECCV). Springer, pp 69\u201385","DOI":"10.1007\/978-3-319-46475-6_5"},{"key":"16913_CR39","doi-asserted-by":"crossref","unstructured":"Mao J, Huang J, Toshev A, Camburu O, Yuille AL, Murphy K (2016) Generation and comprehension of unambiguous object descriptions. In: Proceedings of the IEEE conference on computer vision and pattern recognition (CVPR)","DOI":"10.1109\/CVPR.2016.9"},{"key":"16913_CR40","doi-asserted-by":"crossref","unstructured":"Nagaraja VK, Morariu VI, Davis LS (2016) Modeling context between objects for referring expression understanding. In: European conference on computer vision (ECCV). Springer, pp 792\u2013807","DOI":"10.1007\/978-3-319-46493-0_48"},{"key":"16913_CR41","doi-asserted-by":"crossref","unstructured":"Lin T-Y, Maire M, Belongie S, Hays J, Perona P, Ramanan D, Doll\u00e1r P, Zitnick CL (2014) Microsoft coco: Common objects in context. In: European conference on computer vision. Springer, pp 740\u2013755","DOI":"10.1007\/978-3-319-10602-1_48"},{"key":"16913_CR42","doi-asserted-by":"crossref","unstructured":"Kazemzadeh S, Ordonez V, Matten M, Berg T (2014) Referitgame: Referring to objects in photographs of natural scenes. In: Proceedings of the 2014 conference on empirical methods in natural language processing (EMNLP). pp 787\u2013798","DOI":"10.3115\/v1\/D14-1086"},{"key":"16913_CR43","doi-asserted-by":"crossref","unstructured":"Yang S, Xia M, Li G, Zhou H-Y, Yu Y (2021) Bottom-up shift and reasoning for referring image segmentation. In: Proceedings of the IEEE\/CVF conference on computer vision and pattern recognition (CVPR). pp 11266\u201311275","DOI":"10.1109\/CVPR46437.2021.01111"}],"container-title":["Multimedia Tools and Applications"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11042-023-16913-6.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11042-023-16913-6\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11042-023-16913-6.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2024,4,2]],"date-time":"2024-04-02T13:25:18Z","timestamp":1712064318000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11042-023-16913-6"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2023,9,27]]},"references-count":43,"journal-issue":{"issue":"12","published-online":{"date-parts":[[2024,4]]}},"alternative-id":["16913"],"URL":"https:\/\/doi.org\/10.1007\/s11042-023-16913-6","relation":{},"ISSN":["1573-7721"],"issn-type":[{"type":"electronic","value":"1573-7721"}],"subject":[],"published":{"date-parts":[[2023,9,27]]},"assertion":[{"value":"29 April 2023","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"29 August 2023","order":2,"name":"revised","label":"Revised","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"5 September 2023","order":3,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"27 September 2023","order":4,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no conflicts of interest, and they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflicts of interest"}}]}}