{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2024,9,2]],"date-time":"2024-09-02T14:10:07Z","timestamp":1725286207463},"reference-count":52,"publisher":"Springer Science and Business Media LLC","issue":"29","license":[{"start":{"date-parts":[[2023,12,14]],"date-time":"2023-12-14T00:00:00Z","timestamp":1702512000000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2023,12,14]],"date-time":"2023-12-14T00:00:00Z","timestamp":1702512000000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Multimed Tools Appl"],"DOI":"10.1007\/s11042-023-17696-6","type":"journal-article","created":{"date-parts":[[2023,12,14]],"date-time":"2023-12-14T05:02:19Z","timestamp":1702530139000},"page":"73407-73425","update-policy":"http:\/\/dx.doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":0,"title":["A unified efficient deep image compression framework and its application on human-centric Task"],"prefix":"10.1007","volume":"83","author":[{"given":"Xueyuan","family":"Chen","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Zhihao","family":"Hu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Guo","family":"Lu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Jiaheng","family":"Liu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2023,12,14]]},"reference":[{"key":"17696_CR1","unstructured":"kodak E (2018)kodak lossless true color image suite (photocd pcd0992). http:\/\/r0k.us\/graphics\/kodak\/"},{"key":"17696_CR2","unstructured":"bellard F (2018) bpg image format. http:\/\/bellard.org\/bpg\/, accessed: 30 Oct 2018"},{"key":"17696_CR3","unstructured":"Webp (2018) . https:\/\/developers.google.com\/speed\/webp\/, 30 Oct 2018"},{"key":"17696_CR4","unstructured":"x264 (2018a) the best h.264\/avc encoder. https:\/\/www.videolan.org\/developers\/x264.html, 30 Oct 2018"},{"key":"17696_CR5","unstructured":"x265 (2018b) hevc encoder \/ h.265 video codec. http:\/\/x265.org, 30 Oct 2018"},{"key":"17696_CR6","unstructured":"Agustsson E, Mentzer F, Tschannen M, et al (2017) Soft-to-hard vector quantization for end-to-end learning compressible representations. In: NIPS, pp 1141\u20131151"},{"key":"17696_CR7","doi-asserted-by":"crossref","unstructured":"Agustsson E, Tschannen M, Mentzer F, et al (2018) Generative adversarial networks for extreme learned image compression. arXiv:1804.02958","DOI":"10.1109\/ICCV.2019.00031"},{"key":"17696_CR8","unstructured":"Baig MH, Koltun V, Torresani L (2017) Learning to inpaint for image compression. In: NIPS, pp 1246\u20131255"},{"key":"17696_CR9","unstructured":"Ball\u00e9 J, Laparra V, Simoncelli EP (2015) Density modeling of images using a generalized normalization transformation. arXiv:1511.06281"},{"key":"17696_CR10","unstructured":"Ball\u00e9 J, Laparra V, Simoncelli EP (2017) End-to-end optimized image compression. In: 5th International conference on learning representations, ICLR"},{"key":"17696_CR11","unstructured":"Ball\u00e9 J, Minnen D, Singh S, et al (2018) Variational image compression with a scale hyperprior. In: 6th International conference on learning representations, ICLR"},{"key":"17696_CR12","doi-asserted-by":"crossref","unstructured":"Cao Q, Shen L, Xie W, et al (2018) Vggface2: A dataset for recognising faces across pose and age. In: IEEE International conference on automatic face & gesture recognition. IEEE, pp 67\u201374","DOI":"10.1109\/FG.2018.00020"},{"key":"17696_CR13","doi-asserted-by":"crossref","unstructured":"Chamain LD, Racap\u00e9 F, B\u00e9gaint J, et al (2021) End-to-end optimized image compression for machines, a study. In: 2021 Data compression conference (DCC). IEEE, pp 163\u2013172","DOI":"10.1109\/DCC50243.2021.00024"},{"key":"17696_CR14","unstructured":"Chen T, Liu H, Ma Z, et al (2019) Neural image compression via non-local attention optimization and improved context modeling. arXiv:1910.06244"},{"key":"17696_CR15","doi-asserted-by":"crossref","unstructured":"Cheng Z, Sun H, Takeuchi M, et al (2019) Learning image and video compression through spatial-temporal energy compaction. In: Proceedings of the IEEE conference on computer vision and pattern recognition, CVPR, pp 10,071\u201310,080","DOI":"10.1109\/CVPR.2019.01031"},{"key":"17696_CR16","doi-asserted-by":"crossref","unstructured":"Cheng Z, Sun H, Takeuchi M, et al (2020) Learned image compression with discretized gaussian mixture likelihoods and attention modules. arXiv:2001.01568","DOI":"10.1109\/CVPR42600.2020.00796"},{"key":"17696_CR17","doi-asserted-by":"crossref","unstructured":"Choi Y, El-Khamy M, Lee J (2019) Variable rate deep image compression with a conditional autoencoder. In: Proceedings of the IEEE international conference on computer vision, pp 3146\u20133154","DOI":"10.1109\/ICCV.2019.00324"},{"key":"17696_CR18","doi-asserted-by":"crossref","unstructured":"Deng J, Guo J, Xue N, et al (2019) Arcface: Additive angular margin loss for deep face recognition. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 4690\u20134699","DOI":"10.1109\/CVPR.2019.00482"},{"key":"17696_CR19","doi-asserted-by":"crossref","unstructured":"Djelouah A, Campos J, Schaub-Meyer S, et al (2019) Neural inter-frame compression for video coding. In: Proceedings of the IEEE International conference on computer vision, pp 6421\u20136429","DOI":"10.1109\/ICCV.2019.00652"},{"key":"17696_CR20","doi-asserted-by":"crossref","unstructured":"Duan L, Liu J, Yang W, et al (2020) Video coding for machines: A paradigm of collaborative compression and intelligent analytics. Trans Img Proc 29:8680\u20138695","DOI":"10.1109\/TIP.2020.3016485"},{"key":"17696_CR21","doi-asserted-by":"crossref","unstructured":"Guo Y, Zhang L, Hu Y, et al (2016) Ms-celeb-1m: A dataset and benchmark for large-scale face recognition. In: European conference on computer vision, pp 87\u2013102","DOI":"10.1007\/978-3-319-46487-9_6"},{"key":"17696_CR22","doi-asserted-by":"crossref","unstructured":"Habibian A, Rozendaal Tv, Tomczak JM, et al (2019) Video compression with rate-distortion autoencoders. In: Proceedings of the IEEE international conference on computer vision, pp 7033\u20137042","DOI":"10.1109\/ICCV.2019.00713"},{"key":"17696_CR23","doi-asserted-by":"crossref","unstructured":"Hu J, Shen L, Sun G (2018) Squeeze-and-excitation networks. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 7132\u20137141","DOI":"10.1109\/CVPR.2018.00745"},{"key":"17696_CR24","doi-asserted-by":"crossref","unstructured":"Hu Y, Yang S, Yang W, et al (2020) Towards coding for human and machine vision: A scalable image coding approach. In: 2020 IEEE international conference on multimedia and expo (ICME). IEEE, pp 1\u20136","DOI":"10.1109\/ICME46284.2020.9102750"},{"key":"17696_CR25","doi-asserted-by":"crossref","unstructured":"Johnston N, Vincent D, Minnen D, et al (2018) Improved lossy image compression with priming and spatially adaptive bit rates for recurrent networks. In: CVPR","DOI":"10.1109\/CVPR.2018.00461"},{"key":"17696_CR26","doi-asserted-by":"crossref","unstructured":"Kemelmacher-Shlizerman I, Seitz SM, Miller D, et al (2016) The megaface benchmark: 1 million faces for recognition at scale. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 4873\u20134882","DOI":"10.1109\/CVPR.2016.527"},{"key":"17696_CR27","unstructured":"Kingma DP, Ba J (2014) Adam: A method for stochastic optimization. arXiv:1412.6980"},{"key":"17696_CR28","unstructured":"Lee J, Cho S, Beack SK (2018) Context-adaptive entropy model for end-to-end optimized image compression. arXiv:1809.10452"},{"key":"17696_CR29","unstructured":"Lee J, Cho S, Kim M (2019) A hybrid architecture of jointly learning image compression and quality enhancement with improved entropy minimization. arXiv:1912.12817"},{"key":"17696_CR30","doi-asserted-by":"crossref","unstructured":"Li M, Zuo W, Gu S, et al (2018) Learning convolutional networks for content-weighted image compression. In: CVPR","DOI":"10.1109\/CVPR.2018.00339"},{"key":"17696_CR31","doi-asserted-by":"crossref","unstructured":"Lim B, Son S, Kim H, et al (2017) Enhanced deep residual networks for single image super-resolution. In: Proceedings of the IEEE conference on computer vision and pattern recognition workshops, pp 136\u2013144","DOI":"10.1109\/CVPRW.2017.151"},{"key":"17696_CR32","unstructured":"Liu H, Chen T, Guo P, et al (2019) Non-local attention optimized deep image compression. arXiv:1904.09757"},{"key":"17696_CR33","doi-asserted-by":"crossref","unstructured":"Lu G, Ouyang W, Xu D, et al (2019) DVC: An end-to-end deep video compression framework. In: Proceedings of the IEEE conference on computer vision and pattern recognition,CVPR, pp 11,006\u201311,015","DOI":"10.1109\/CVPR.2019.01126"},{"key":"17696_CR34","doi-asserted-by":"crossref","unstructured":"Mentzer F, Agustsson E, Tschannen M, et al (2018) Conditional probability models for deep image compression. In: CVPR, 2, p\u00a03","DOI":"10.1109\/CVPR.2018.00462"},{"key":"17696_CR35","unstructured":"Minnen D, Ball\u00e9 J, Toderici GD (2018) Joint autoregressive and hierarchical priors for learned image compression. In: Advances in neural information processing systems, pp 10,771\u201310,780"},{"key":"17696_CR36","unstructured":"Nair V, Hinton GE (2010) Rectified linear units improve restricted boltzmann machines. In: Proceedings of the 27th international conference on machine learning (ICML-10), pp 807\u2013814"},{"key":"17696_CR37","unstructured":"Paszke A, Gross S, Massa F, et al (2019) Pytorch: An imperative style, high-performance deep learning library. In: Advances in neural information processing systems 32. Curran Associates, Inc., p 8024\u20138035"},{"key":"17696_CR38","doi-asserted-by":"crossref","unstructured":"Ranjan A, Black MJ (2017) Optical flow estimation using a spatial pyramid network. In: The IEEE Conference on computer vision and pattern recognition (CVPR)","DOI":"10.1109\/CVPR.2017.291"},{"key":"17696_CR39","unstructured":"Rippel O, Bourdev L (2017) Real-time adaptive image compression. In: ICML"},{"key":"17696_CR40","doi-asserted-by":"crossref","unstructured":"Sandler M, Howard A, Zhu M, et\u00a0al (2018) Mobilenetv2: Inverted residuals and linear bottlenecks. In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, pp 4510\u20134520","DOI":"10.1109\/CVPR.2018.00474"},{"issue":"5","key":"17696_CR41","doi-asserted-by":"publisher","first-page":"36","DOI":"10.1109\/79.952804","volume":"18","author":"A Skodras","year":"2001","unstructured":"Skodras A, Christopoulos C, Ebrahimi T (2001) The jpeg 2000 still image compression standard. IEEE Signal Process Mag 18(5):36\u201358","journal-title":"IEEE Signal Process Mag"},{"issue":"12","key":"17696_CR42","first-page":"1649","volume":"22","author":"GJ Sullivan","year":"2012","unstructured":"Sullivan GJ, Ohm JR, Han WJ et al (2012) Overview of the high efficiency video coding(hevc) standard. TCSVT 22(12):1649\u20131668","journal-title":"TCSVT"},{"key":"17696_CR43","unstructured":"Theis L, Shi W, Cunningham A, et\u00a0al (2017) Lossy image compression with compressive autoencoders. In: 5th International conference on learning representations, ICLR"},{"key":"17696_CR44","doi-asserted-by":"crossref","unstructured":"Toderici G, O\u2019Malley SM, Hwang SJ, et\u00a0al (2016) Variable rate image compression with recurrent neural networks. In: 4th International conference on learning representations, ICLR","DOI":"10.1109\/CVPR.2017.577"},{"key":"17696_CR45","doi-asserted-by":"crossref","unstructured":"Toderici G, Vincent D, Johnston N, et al (2017) Full resolution image compression with recurrent neural networks. In: CVPR, pp 5435\u20135443","DOI":"10.1109\/CVPR.2017.577"},{"key":"17696_CR46","doi-asserted-by":"crossref","unstructured":"Wallace GK (1992) The jpeg still picture compression standard. IEEE Transactions on Consumer Electronics 38(1):xviii\u2013xxxiv","DOI":"10.1109\/30.125072"},{"key":"17696_CR47","doi-asserted-by":"crossref","unstructured":"Wang X, Girshick R, Gupta A, et\u00a0al (2018) Non-local neural networks. In: Proceedings of the IEEE conference on computer vision and pattern recognition, pp 7794\u20137803","DOI":"10.1109\/CVPR.2018.00813"},{"key":"17696_CR48","unstructured":"Wang Z, Simoncelli E, Bovik A, et al (2003) Multi-scale structural similarity for image quality assessment. In: ASILOMAR CONFERENCE ON SIGNALS SYSTEMS AND COMPUTERS, IEEE; 1998, pp 1398\u20131402"},{"issue":"6","key":"17696_CR49","doi-asserted-by":"publisher","first-page":"520","DOI":"10.1145\/214762.214771","volume":"30","author":"IH Witten","year":"1987","unstructured":"Witten IH, Neal RM, Cleary JG (1987) Arithmetic coding for data compression. Communications of the ACM 30(6):520\u2013540","journal-title":"Communications of the ACM"},{"key":"17696_CR50","doi-asserted-by":"crossref","unstructured":"Wu CY, Singhal N, Krahenbuhl P (2018) Video compression through image interpolation. In: ECCV","DOI":"10.1007\/978-3-030-01237-3_26"},{"issue":"8","key":"17696_CR51","doi-asserted-by":"publisher","first-page":"1106","DOI":"10.1007\/s11263-018-01144-2","volume":"127","author":"T Xue","year":"2019","unstructured":"Xue T, Chen B, Wu J et al (2019) Video enhancement with task-oriented flow. International Journal of Computer Vision, IJCV 127(8):1106\u20131125","journal-title":"International Journal of Computer Vision, IJCV"},{"key":"17696_CR52","doi-asserted-by":"publisher","first-page":"58","DOI":"10.1016\/j.neucom.2022.08.048","volume":"508","author":"F Yang","year":"2022","unstructured":"Yang F, Wang Y, Herranz L et al (2022) A novel framework for image-to-image translation and image compression. Neurocomputing 508:58\u201370","journal-title":"Neurocomputing"}],"container-title":["Multimedia Tools and Applications"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11042-023-17696-6.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s11042-023-17696-6\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s11042-023-17696-6.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2024,9,2]],"date-time":"2024-09-02T13:18:39Z","timestamp":1725283119000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s11042-023-17696-6"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2023,12,14]]},"references-count":52,"journal-issue":{"issue":"29","published-online":{"date-parts":[[2024,9]]}},"alternative-id":["17696"],"URL":"https:\/\/doi.org\/10.1007\/s11042-023-17696-6","relation":{},"ISSN":["1573-7721"],"issn-type":[{"type":"electronic","value":"1573-7721"}],"subject":[],"published":{"date-parts":[[2023,12,14]]},"assertion":[{"value":"13 June 2023","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"4 September 2023","order":2,"name":"revised","label":"Revised","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"21 November 2023","order":3,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"14 December 2023","order":4,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflicts of interest"}}]}}