{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,7,13]],"date-time":"2026-07-13T12:13:11Z","timestamp":1783944791202,"version":"3.55.0"},"reference-count":184,"publisher":"Elsevier BV","license":[{"start":{"date-parts":[[2026,9,1]],"date-time":"2026-09-01T00:00:00Z","timestamp":1788220800000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.elsevier.com\/tdm\/userlicense\/1.0\/"},{"start":{"date-parts":[[2026,9,1]],"date-time":"2026-09-01T00:00:00Z","timestamp":1788220800000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.elsevier.com\/legal\/tdmrep-license"},{"start":{"date-parts":[[2026,9,1]],"date-time":"2026-09-01T00:00:00Z","timestamp":1788220800000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-017"},{"start":{"date-parts":[[2026,9,1]],"date-time":"2026-09-01T00:00:00Z","timestamp":1788220800000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-037"},{"start":{"date-parts":[[2026,9,1]],"date-time":"2026-09-01T00:00:00Z","timestamp":1788220800000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-012"},{"start":{"date-parts":[[2026,9,1]],"date-time":"2026-09-01T00:00:00Z","timestamp":1788220800000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-029"},{"start":{"date-parts":[[2026,9,1]],"date-time":"2026-09-01T00:00:00Z","timestamp":1788220800000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-004"}],"funder":[{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"publisher","award":["62372018"],"award-info":[{"award-number":["62372018"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"publisher"}]},{"DOI":"10.13039\/501100001809","name":"National Natural Science Foundation of China","doi-asserted-by":"publisher","award":["62272016"],"award-info":[{"award-number":["62272016"]}],"id":[{"id":"10.13039\/501100001809","id-type":"DOI","asserted-by":"publisher"}]}],"content-domain":{"domain":["elsevier.com","sciencedirect.com"],"crossmark-restriction":true},"short-container-title":["Neurocomputing"],"published-print":{"date-parts":[[2026,9]]},"DOI":"10.1016\/j.neucom.2026.133839","type":"journal-article","created":{"date-parts":[[2026,5,12]],"date-time":"2026-05-12T09:57:18Z","timestamp":1778579838000},"page":"133839","update-policy":"https:\/\/doi.org\/10.1016\/elsevier_cm_policy","source":"Crossref","is-referenced-by-count":0,"special_numbering":"C","title":["End-to-end learned video compression: A comprehensive review"],"prefix":"10.1016","volume":"694","author":[{"ORCID":"https:\/\/orcid.org\/0000-0002-8402-2981","authenticated-orcid":false,"given":"Huanjie","family":"He","sequence":"first","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Yunhui","family":"Shi","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"ORCID":"https:\/\/orcid.org\/0000-0001-5437-3150","authenticated-orcid":false,"given":"Jin","family":"Wang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Jiuzhen","family":"Liang","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Nam","family":"Ling","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]},{"given":"Baocai","family":"Yin","sequence":"additional","affiliation":[],"role":[{"vocabulary":"crossref","role":"author"}]}],"member":"78","reference":[{"key":"10.1016\/j.neucom.2026.133839_bib0005","series-title":"Advances in Neural Information Processing Systems","article-title":"Soft-to-hard vector quantization for end-to-end learning compressible representations","author":"Agustsson","year":"2017"},{"key":"10.1016\/j.neucom.2026.133839_bib0010","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition","first-page":"8503","article-title":"Scale-space flow for end-to-end optimized video compression","author":"Agustsson","year":"2020"},{"key":"10.1016\/j.neucom.2026.133839_bib0015","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"10249","article-title":"Hierarchical B-Frame video coding using two-layer CANF without motion coding","author":"Alexandre","year":"2023"},{"key":"10.1016\/j.neucom.2026.133839_bib0020","series-title":"2021 IEEE International Conference on Image Processing (ICIP)","first-page":"2124","article-title":"Deep video compression for interframe coding","author":"Alexandre","year":"2021"},{"key":"10.1016\/j.neucom.2026.133839_bib0025","series-title":"2023 IEEE International Conference on Image Processing (ICIP)","first-page":"41","article-title":"PS-NeRV: patch-wise stylized neural representations for videos","author":"Bai","year":"2023"},{"key":"10.1016\/j.neucom.2026.133839_bib0030","series-title":"5th International Conference on Learning Representations, ICLR 2017, Toulon, France","article-title":"End-to-end optimized image compression","author":"Ball\u00e9","year":"2017"},{"key":"10.1016\/j.neucom.2026.133839_bib0035","series-title":"2016 Picture Coding Symposium (PCS)","first-page":"1","article-title":"End-to-end optimization of nonlinear transform codes for perceptual quality","author":"Ball\u00e9","year":"2016"},{"key":"10.1016\/j.neucom.2026.133839_bib0040","author":"Ball\u00e9"},{"key":"10.1016\/j.neucom.2026.133839_bib0045","series-title":"Rate-distortion theory","author":"Berger","year":"2003"},{"key":"10.1016\/j.neucom.2026.133839_bib0050","series-title":"Proceedings of the 36th International Conference on Machine Learning","first-page":"675","article-title":"Rethinking lossy compression: the rate-distortion-perception tradeoff","author":"Blau","year":"2019"},{"key":"10.1016\/j.neucom.2026.133839_bib0055","series-title":"2022 Picture Coding Symposium (PCS)","first-page":"289","article-title":"On benefits and challenges of conditional interframe video coding in light of information theory","author":"Brand","year":"2022"},{"key":"10.1016\/j.neucom.2026.133839_bib0060","series-title":"2022 IEEE International Conference on Image Processing (ICIP)","first-page":"1266","article-title":"P-Frame coding with generalized difference: a novel conditional coding approach","author":"Brand","year":"2022"},{"key":"10.1016\/j.neucom.2026.133839_bib0065","doi-asserted-by":"crossref","first-page":"6445","DOI":"10.1109\/TCSVT.2024.3359948","article-title":"Conditional residual coding: a remedy for bottleneck problems in conditional inter frame coding","volume":"34","author":"Brand","year":"2024","journal-title":"IEEE Trans. Circuits Syst. Video Technol."},{"key":"10.1016\/j.neucom.2026.133839_bib0070","doi-asserted-by":"crossref","first-page":"3736","DOI":"10.1109\/TCSVT.2021.3101953","article-title":"Overview of the versatile video coding (VVC) standard and its applications","volume":"31","author":"Bross","year":"2021","journal-title":"IEEE Trans. Circuits Syst. Video Technol."},{"key":"10.1016\/j.neucom.2026.133839_bib0075","series-title":"Proceedings of the AAAI Conference on Artificial Intelligence","first-page":"973","article-title":"Understanding deformable alignment in video super-resolution","author":"Chan","year":"2021"},{"key":"10.1016\/j.neucom.2026.133839_bib0080","doi-asserted-by":"crossref","first-page":"2910","DOI":"10.1109\/TIP.2025.3563762","article-title":"Interactive face video coding: a generative compression framework","volume":"34","author":"Chen","year":"2025","journal-title":"IEEE Trans. Image Process."},{"key":"10.1016\/j.neucom.2026.133839_bib0085","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"10270","article-title":"HNeRV: a hybrid neural representation for videos","author":"Chen","year":"2023"},{"key":"10.1016\/j.neucom.2026.133839_bib0090","series-title":"Advances in Neural Information Processing Systems","first-page":"21557","article-title":"NeRV: neural representations for videos","author":"Chen","year":"2021"},{"key":"10.1016\/j.neucom.2026.133839_bib0095","series-title":"2021 Picture Coding Symposium (PCS)","first-page":"1","article-title":"MOVI-codec: deep video compression without motion","author":"Chen","year":"2021"},{"key":"10.1016\/j.neucom.2026.133839_bib0100","doi-asserted-by":"crossref","first-page":"2908","DOI":"10.1109\/TCSVT.2023.3301016","article-title":"B-CANF: adaptive B-Frame coding with conditional augmented normalizing flows","volume":"34","author":"Chen","year":"2024","journal-title":"IEEE Trans. Circuits Syst. Video Technol."},{"key":"10.1016\/j.neucom.2026.133839_bib0105","doi-asserted-by":"crossref","first-page":"11980","DOI":"10.1109\/TCSVT.2024.3427426","article-title":"MaskCRT: masked conditional residual transformer for learned video compression","volume":"34","author":"Chen","year":"2024","journal-title":"IEEE Trans. Circuits Syst. Video Technol."},{"key":"10.1016\/j.neucom.2026.133839_bib0110","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","article-title":"Learned image compression with discretized Gaussian mixture likelihoods and attention modules","author":"Cheng","year":"2020"},{"key":"10.1016\/j.neucom.2026.133839_bib0115","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"10532","article-title":"Asymmetric gained deep image compression with continuous rate adaptation","author":"Cui","year":"2021"},{"key":"10.1016\/j.neucom.2026.133839_bib0120","series-title":"Proceedings of the IEEE International Conference on Computer Vision (ICCV)","article-title":"Deformable convolutional networks","author":"Dai","year":"2017"},{"key":"10.1016\/j.neucom.2026.133839_bib0125","series-title":"Proceedings of the IEEE\/CVF International Conference on Computer Vision (ICCV)","article-title":"Neural inter-frame compression for video coding","author":"Djelouah","year":"2019"},{"key":"10.1016\/j.neucom.2026.133839_bib0130","author":"Dong"},{"key":"10.1016\/j.neucom.2026.133839_bib0135","doi-asserted-by":"crossref","first-page":"209","DOI":"10.1109\/JETCAS.2024.3387301","article-title":"CGVC-T: contextual generative video compression with transformers","volume":"14","author":"Du","year":"2024","journal-title":"IEEE J. Emerg. Sel. Top. Circuits Syst."},{"key":"10.1016\/j.neucom.2026.133839_bib0140","series-title":"Big Data IV: Learning, Analytics, and Applications","article-title":"A generative adversarial network for video compression","author":"Du","year":"2022"},{"key":"10.1016\/j.neucom.2026.133839_bib0145","series-title":"2022 Picture Coding Symposium (PCS)","first-page":"349","article-title":"Generative video compression with a transformer-based discriminator","author":"Du","year":"2022"},{"key":"10.1016\/j.neucom.2026.133839_bib0160","series-title":"Proceedings of the IEEE\/CVF International Conference on Computer Vision (ICCV)","first-page":"17356","article-title":"GIViC: generative implicit video compression","author":"Gao","year":"2025"},{"key":"10.1016\/j.neucom.2026.133839_bib0165","doi-asserted-by":"crossref","first-page":"3041","DOI":"10.1109\/TIP.2025.3567830","article-title":"Approximately invertible neural network for learned image compression","volume":"34","author":"Gao","year":"2025","journal-title":"IEEE Trans. Image Process."},{"key":"10.1016\/j.neucom.2026.133839_bib0170","series-title":"2024 IEEE International Symposium on Circuits and Systems (ISCAS)","first-page":"1","article-title":"Conditional variational autoencoders for hierarchical B-frame coding","author":"Gao","year":"2024"},{"key":"10.1016\/j.neucom.2026.133839_bib0175","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"18497","article-title":"Video compression with entropy-constrained neural representations","author":"Gomes","year":"2023"},{"key":"10.1016\/j.neucom.2026.133839_bib0180","series-title":"Advances in Neural Information Processing Systems","article-title":"Generative adversarial nets","author":"Goodfellow","year":"2014"},{"key":"10.1016\/j.neucom.2026.133839_bib0185","doi-asserted-by":"crossref","first-page":"673","DOI":"10.1109\/LSP.2023.3277343","article-title":"Enhanced motion compensation for deep video compression","volume":"30","author":"Guo","year":"2023","journal-title":"IEEE Signal Process. Lett."},{"key":"10.1016\/j.neucom.2026.133839_bib0190","doi-asserted-by":"crossref","first-page":"3814","DOI":"10.1109\/TMM.2023.3316429","article-title":"Enhanced context mining and filtering for learned video compression","volume":"26","author":"Guo","year":"2024","journal-title":"IEEE Trans. Multimed."},{"key":"10.1016\/j.neucom.2026.133839_bib0195","doi-asserted-by":"crossref","first-page":"517","DOI":"10.1109\/TBC.2025.3541869","article-title":"Exploring invertible encoding for deep video compression","volume":"71","author":"Guo","year":"2025","journal-title":"IEEE Trans. Broadcast."},{"key":"10.1016\/j.neucom.2026.133839_bib0200","doi-asserted-by":"crossref","first-page":"3567","DOI":"10.1109\/TIP.2023.3287495","article-title":"Learning cross-scale weighted prediction for efficient neural video compression","volume":"32","author":"Guo","year":"2023","journal-title":"IEEE Trans. Image Process."},{"key":"10.1016\/j.neucom.2026.133839_bib0205","series-title":"Proceedings of the IEEE\/CVF International Conference on Computer Vision (ICCV)","article-title":"Video compression with rate-distortion autoencoders","author":"Habibian","year":"2019"},{"key":"10.1016\/j.neucom.2026.133839_bib0210","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"6132","article-title":"Towards scalable neural representation for diverse videos","author":"He","year":"2023"},{"key":"10.1016\/j.neucom.2026.133839_bib0215","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"5718","article-title":"ELIC: efficient learned image compression with unevenly grouped space-channel contextual adaptive coding","author":"He","year":"2022"},{"key":"10.1016\/j.neucom.2026.133839_bib0220","series-title":"Advances in Neural Information Processing Systems","article-title":"GANs trained by a two time-scale update rule converge to a local Nash equilibrium","author":"Heusel","year":"2017"},{"key":"10.1016\/j.neucom.2026.133839_bib0225","series-title":"Computer Vision \u2013 ECCV 2022","first-page":"207","article-title":"CANF-VC: conditional augmented normalizing flows for video compression","author":"Ho","year":"2022"},{"key":"10.1016\/j.neucom.2026.133839_bib0230","series-title":"Computer Vision \u2013 ECCV 2020","first-page":"193","article-title":"Improving deep video compression by resolution-adaptive flow coding","author":"Hu","year":"2020"},{"key":"10.1016\/j.neucom.2026.133839_bib0235","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"5921","article-title":"Coarse-to-fine deep video coding with hyperprior-guided mode prediction","author":"Hu","year":"2022"},{"key":"10.1016\/j.neucom.2026.133839_bib0240","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"1502","article-title":"FVC: a new framework towards deep video compression in feature space","author":"Hu","year":"2021"},{"key":"10.1016\/j.neucom.2026.133839_bib0245","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"14358","article-title":"Complexity-guided slimmable decoder for efficient deep video compression","author":"Hu","year":"2023"},{"key":"10.1016\/j.neucom.2026.133839_bib0250","series-title":"2025 International Symposium on Machine Learning and Media Computing (MLMC)","first-page":"1","article-title":"Referenced frame generation network for learned video compression","author":"Huang","year":"2025"},{"key":"10.1016\/j.neucom.2026.133839_bib0255","series-title":"2025 International Symposium on Machine Learning and Media Computing (MLMC)","first-page":"1","article-title":"Multi-frame context generation for neural video compression","author":"Huang","year":"2025"},{"key":"10.1016\/j.neucom.2026.133839_bib0260","series-title":"Proceedings of the IEEE International Conference on Computer Vision (ICCV)","article-title":"Arbitrary style transfer in real-time with adaptive instance normalization","author":"Huang","year":"2017"},{"key":"10.1016\/j.neucom.2026.133839_bib0265","series-title":"Computer Vision \u2013 ECCV 2022","first-page":"624","article-title":"Real-time intermediate flow estimation for video frame interpolation","author":"Huang","year":"2022"},{"key":"10.1016\/j.neucom.2026.133839_bib0270","series-title":"Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (CVPR)","article-title":"LiteFlowNet: a lightweight convolutional neural network for optical flow estimation","author":"Hui","year":"2018"},{"key":"10.1016\/j.neucom.2026.133839_bib0275","doi-asserted-by":"crossref","DOI":"10.1145\/3072959.3073659","article-title":"Globally and locally consistent image completion","volume":"36","author":"Iizuka","year":"2017","journal-title":"ACM Trans. Graph."},{"key":"10.1016\/j.neucom.2026.133839_bib0280","doi-asserted-by":"crossref","first-page":"3096","DOI":"10.1109\/TCSVT.2023.3310188","article-title":"MPAI-EEV: standardization efforts of artificial intelligence based end-to-end video coding","volume":"34","author":"Jia","year":"2024","journal-title":"IEEE Trans. Circuits Syst. Video Technol."},{"key":"10.1016\/j.neucom.2026.133839_bib0285","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"12543","article-title":"Towards practical real-time neural video compression","author":"Jia","year":"2025"},{"key":"10.1016\/j.neucom.2026.133839_bib0290","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"26088","article-title":"Generative latent coding for ultra-low bitrate image compression","author":"Jia","year":"2024"},{"key":"10.1016\/j.neucom.2026.133839_bib0295","series-title":"Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (CVPR)","article-title":"Super SloMo: high quality estimation of multiple intermediate frames for video interpolation","author":"Jiang","year":"2018"},{"key":"10.1016\/j.neucom.2026.133839_bib0300","doi-asserted-by":"crossref","first-page":"1800","DOI":"10.1109\/LSP.2024.3411917","article-title":"OMR-NET: a two-stage octave multi-scale residual network for screen content image compression","volume":"31","author":"Jiang","year":"2024","journal-title":"IEEE Signal Process. Lett."},{"key":"10.1016\/j.neucom.2026.133839_bib0305","doi-asserted-by":"crossref","first-page":"3290","DOI":"10.1109\/LSP.2025.3596872","article-title":"OMR-net+: a frequency-aware feature refinement and entropy modeling method for efficient screen content image compression","volume":"32","author":"Jiang","year":"2025","journal-title":"IEEE Signal Process. Lett."},{"key":"10.1016\/j.neucom.2026.133839_bib0310","series-title":"2023 IEEE International Conference on Image Processing (ICIP)","first-page":"2975","article-title":"Fourier series and laplacian noise-based quantization error compensation for end-to-end learning-based image compression","author":"Jiang","year":"2023"},{"key":"10.1016\/j.neucom.2026.133839_bib0315","doi-asserted-by":"crossref","first-page":"1034","DOI":"10.1109\/TBC.2025.3609052","article-title":"FD-LSCIC: frequency decomposition-based learned screen content image compression","volume":"71","author":"Jiang","year":"2025","journal-title":"IEEE Trans. Broadcast."},{"key":"10.1016\/j.neucom.2026.133839_bib0320","series-title":"ICASSP 2024 - 2024 IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP)","first-page":"2955","article-title":"LVC-LGMC: joint local and global motion compensation for learned video compression","author":"Jiang","year":"2024"},{"key":"10.1016\/j.neucom.2026.133839_bib0325","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"7331","article-title":"ECVC: exploiting non-local correlations in multiple frames for contextual video compression","author":"Jiang","year":"2025"},{"key":"10.1016\/j.neucom.2026.133839_bib0330","doi-asserted-by":"crossref","first-page":"3188","DOI":"10.1109\/TIP.2023.3276333","article-title":"Learned video compression with efficient temporal context learning","volume":"32","author":"Jin","year":"2023","journal-title":"IEEE Trans. Image Process."},{"key":"10.1016\/j.neucom.2026.133839_bib0335","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"9347","article-title":"C3: high-performance and low-complexity neural compression from a single image or video","author":"Kim","year":"2024"},{"key":"10.1016\/j.neucom.2026.133839_bib0340","series-title":"Advances in Neural Information Processing Systems","first-page":"12718","article-title":"Scalable neural video representations with learnable positional features","author":"Kim","year":"2022"},{"key":"10.1016\/j.neucom.2026.133839_bib0345","series-title":"Advances in Neural Information Processing Systems","first-page":"72692","article-title":"HiNeRV: video compression with hierarchical encoding-based neural representation","author":"Kwan","year":"2023"},{"key":"10.1016\/j.neucom.2026.133839_bib0350","series-title":"2022 IEEE International Conference on Image Processing (ICIP)","first-page":"316","article-title":"AIVC: artificial intelligence based video codec","author":"Ladune","year":"2022"},{"key":"10.1016\/j.neucom.2026.133839_bib0355","series-title":"2020 IEEE 30th International Workshop on Machine Learning for Signal Processing (MLSP)","first-page":"1","article-title":"ModeNet: mode selection network for learned video coding","author":"Ladune","year":"2020"},{"key":"10.1016\/j.neucom.2026.133839_bib0360","series-title":"2020 IEEE 22nd International Workshop on Multimedia Signal Processing (MMSP)","first-page":"1","article-title":"Optical flow and mode selection for learning-based video coding","author":"Ladune","year":"2020"},{"key":"10.1016\/j.neucom.2026.133839_bib0365","series-title":"Proceedings of the IEEE\/CVF International Conference on Computer Vision (ICCV)","first-page":"13515","article-title":"COOL-CHIC: coordinate-based low complexity hierarchical image codec","author":"Ladune","year":"2023"},{"key":"10.1016\/j.neucom.2026.133839_bib0370","series-title":"Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (CVPR)","article-title":"Photo-realistic single image super-resolution using a generative adversarial network","author":"Ledig","year":"2017"},{"key":"10.1016\/j.neucom.2026.133839_bib0375","series-title":"Proceedings of the 31st ACM International Conference on Multimedia","first-page":"7859","article-title":"FFNeRV: flow-guided frame-wise neural representations for videos","author":"Lee","year":"2023"},{"key":"10.1016\/j.neucom.2026.133839_bib0380","series-title":"2023 IEEE 25th International Workshop on Multimedia Signal Processing (MMSP)","first-page":"1","article-title":"Low-complexity overfitted neural image codec","author":"Leguay","year":"2023"},{"key":"10.1016\/j.neucom.2026.133839_bib0385","series-title":"2024 Data Compression Conference (DCC)","first-page":"23","article-title":"Cool-chic video: learned video coding with 800 parameters","author":"Leguay","year":"2024"},{"key":"10.1016\/j.neucom.2026.133839_bib0390","series-title":"2024 16th International Conference on Wireless Communications and Signal Processing (WCSP)","first-page":"1449","article-title":"Extreme video compression with prediction using pre-trained diffusion models","author":"Li","year":"2024"},{"key":"10.1016\/j.neucom.2026.133839_bib0395","series-title":"ICASSP 2022 - 2022 IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP)","first-page":"1855","article-title":"Jmpnet: joint motion prediction for learning-based video compression","author":"Li","year":"2022"},{"key":"10.1016\/j.neucom.2026.133839_bib0400","series-title":"Advances in Neural Information Processing Systems","first-page":"18114","article-title":"Deep contextual video compression","author":"Li","year":"2021"},{"key":"10.1016\/j.neucom.2026.133839_bib0405","series-title":"Proceedings of the 30th ACM International Conference on Multimedia","first-page":"1503","article-title":"Hybrid spatial-temporal entropy modelling for neural video compression","author":"Li","year":"2022"},{"key":"10.1016\/j.neucom.2026.133839_bib0410","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"22616","article-title":"Neural video compression with diverse contexts","author":"Li","year":"2023"},{"key":"10.1016\/j.neucom.2026.133839_bib0415","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"26099","article-title":"Neural video compression with feature modulation","author":"Li","year":"2024"},{"key":"10.1016\/j.neucom.2026.133839_bib0420","series-title":"Proceedings of the 31st ACM International Conference on Multimedia","first-page":"8057","article-title":"High visual-fidelity learned video compression","author":"Li","year":"2023"},{"key":"10.1016\/j.neucom.2026.133839_bib0425","doi-asserted-by":"crossref","first-page":"6576","DOI":"10.1109\/TII.2022.3204681","article-title":"Learning-based video compression framework with implicit spatial transform for applications in the internet of things","volume":"19","author":"Li","year":"2023","journal-title":"IEEE Trans. Ind. Inform."},{"key":"10.1016\/j.neucom.2026.133839_bib0430","series-title":"Computer Vision \u2013 ECCV 2022","first-page":"267","article-title":"E-NeRV: expedite neural video representation with disentangled spatial-temporal context","author":"Li","year":"2022"},{"key":"10.1016\/j.neucom.2026.133839_bib0435","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"3546","article-title":"M-LVC: multiple frames prediction for learned video compression","author":"Lin","year":"2020"},{"key":"10.1016\/j.neucom.2026.133839_bib0440","doi-asserted-by":"crossref","first-page":"3502","DOI":"10.1109\/TCSVT.2022.3233221","article-title":"DMVC: decomposed motion modeling for learned video compression","volume":"33","author":"Lin","year":"2023","journal-title":"IEEE Trans. Circuits Syst. Video Technol."},{"key":"10.1016\/j.neucom.2026.133839_bib0445","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"701","article-title":"Deep learning in latent space for video prediction and compression","author":"Liu","year":"2021"},{"key":"10.1016\/j.neucom.2026.133839_bib0450","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"18487","article-title":"MMVC: learned multi-mode video compression with block-based prediction mode selection and density-adaptive entropy coding","author":"Liu","year":"2023"},{"key":"10.1016\/j.neucom.2026.133839_bib0455","series-title":"2022 IEEE International Conference on Image Processing (ICIP)","first-page":"1321","article-title":"Learned video compression with residual prediction and feature-aided loop filter","author":"Liu","year":"2022"},{"key":"10.1016\/j.neucom.2026.133839_bib0460","article-title":"Deep learning-based video coding: a review and a case study","volume":"53","author":"Liu","year":"2020","journal-title":"ACM Comput. Surv."},{"key":"10.1016\/j.neucom.2026.133839_bib0465","doi-asserted-by":"crossref","first-page":"5650","DOI":"10.1109\/TCSVT.2022.3150014","article-title":"End-to-end neural video coding using a compound spatiotemporal representation","volume":"32","author":"Liu","year":"2022","journal-title":"IEEE Trans. Circuits Syst. Video Technol."},{"key":"10.1016\/j.neucom.2026.133839_bib0470","doi-asserted-by":"crossref","first-page":"3182","DOI":"10.1109\/TCSVT.2020.3035680","article-title":"Neural video coding using multiscale motion compensation and spatiotemporal context model","volume":"31","author":"Liu","year":"2021","journal-title":"IEEE Trans. Circuits Syst. Video Technol."},{"key":"10.1016\/j.neucom.2026.133839_bib0475","series-title":"Proceedings of the AAAI Conference on Artificial Intelligence","first-page":"11580","article-title":"Learned video compression via joint spatial-temporal correlation exploration","author":"Liu","year":"2020"},{"key":"10.1016\/j.neucom.2026.133839_bib0480","series-title":"European Conference on Computer Vision","first-page":"453","article-title":"Conditional entropy coding for efficient video compression","author":"Liu","year":"2020"},{"key":"10.1016\/j.neucom.2026.133839_bib0485","author":"Liu"},{"key":"10.1016\/j.neucom.2026.133839_bib0490","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"11976","article-title":"A ConvNet for the 2020s","author":"Liu","year":"2022"},{"key":"10.1016\/j.neucom.2026.133839_bib0495","series-title":"Advances in Neural Information Processing Systems","article-title":"Deep generative video compression","author":"Lombardo","year":"2019"},{"key":"10.1016\/j.neucom.2026.133839_bib0500","series-title":"Computer Vision\u2013ECCV 2020: 16th European Conference, Glasgow, UK, August 23\u201328, 2020, Proceedings, Part II 16","first-page":"456","article-title":"Content adaptive and error propagation aware deep video compression","author":"Lu","year":"2020"},{"key":"10.1016\/j.neucom.2026.133839_bib0505","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","article-title":"DVC: an end-to-end deep video compression framework","author":"Lu","year":"2019"},{"key":"10.1016\/j.neucom.2026.133839_bib0510","doi-asserted-by":"crossref","first-page":"3292","DOI":"10.1109\/TPAMI.2020.2988453","article-title":"An end-to-end learning framework for video compression","volume":"43","author":"Lu","year":"2021","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"10.1016\/j.neucom.2026.133839_bib0515","series-title":"Proceedings of the AAAI Conference on Artificial Intelligence","first-page":"8859","article-title":"Deep hierarchical video compression","author":"Lu","year":"2024"},{"key":"10.1016\/j.neucom.2026.133839_bib0520","doi-asserted-by":"crossref","DOI":"10.1145\/3761815","article-title":"Diffusion-based perceptual neural video compression with temporal diffusion information reuse","volume":"21","author":"Ma","year":"2025","journal-title":"ACM Trans. Multimedia Comput. Commun. Appl."},{"key":"10.1016\/j.neucom.2026.133839_bib0525","author":"Ma"},{"key":"10.1016\/j.neucom.2026.133839_bib0530","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"10738","article-title":"Motion-adjustable neural implicit video representation","author":"Mai","year":"2022"},{"key":"10.1016\/j.neucom.2026.133839_bib0535","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"14378","article-title":"NIRVANA: neural implicit representations of videos with adaptive networks and autoregressive patch-wise modeling","author":"Maiya","year":"2023"},{"key":"10.1016\/j.neucom.2026.133839_bib0540","author":"Mao"},{"key":"10.1016\/j.neucom.2026.133839_bib0545","series-title":"European Conference on Computer Vision","first-page":"562","article-title":"Neural video compression using GANs for detail synthesis and propagation","author":"Mentzer","year":"2022"},{"key":"10.1016\/j.neucom.2026.133839_bib0550","author":"Mentzer"},{"key":"10.1016\/j.neucom.2026.133839_bib0555","doi-asserted-by":"crossref","first-page":"99","DOI":"10.1145\/3503250","article-title":"NeRF: representing scenes as neural radiance fields for view synthesis","volume":"65","author":"Mildenhall","year":"2021","journal-title":"Commun. ACM"},{"key":"10.1016\/j.neucom.2026.133839_bib0560","series-title":"Advances in Neural Information Processing Systems","article-title":"Joint autoregressive and hierarchical priors for learned image compression","author":"Minnen","year":"2018"},{"key":"10.1016\/j.neucom.2026.133839_bib0565","series-title":"2020 IEEE International Conference on Image Processing (ICIP)","first-page":"3339","article-title":"Channel-wise autoregressive entropy models for learned image compression","author":"Minnen","year":"2020"},{"key":"10.1016\/j.neucom.2026.133839_bib0570","doi-asserted-by":"crossref","first-page":"23","DOI":"10.1109\/79.733495","article-title":"Rate-distortion methods for image and video compression","volume":"15","author":"Ortega","year":"1998","journal-title":"IEEE Signal Process. Mag."},{"key":"10.1016\/j.neucom.2026.133839_bib0575","series-title":"Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (CVPR)","article-title":"Context encoders: feature learning by inpainting","author":"Pathak","year":"2016"},{"key":"10.1016\/j.neucom.2026.133839_bib0580","series-title":"2020 IEEE Workshop on Signal Processing Systems (SIPS)","first-page":"1","article-title":"End-to-end learning of video compression using spatio-temporal autoencoders","author":"Pessoa","year":"2020"},{"key":"10.1016\/j.neucom.2026.133839_bib0585","series-title":"Proceedings of the IEEE\/CVF International Conference on Computer Vision (ICCV)","first-page":"6680","article-title":"Extending neural P-Frame codecs for B-Frame coding","author":"Pourreza","year":"2021"},{"key":"10.1016\/j.neucom.2026.133839_bib0590","doi-asserted-by":"crossref","first-page":"10500","DOI":"10.1109\/TCSVT.2025.3571944","article-title":"Generative latent coding for ultra-low bitrate image and video compression","volume":"35","author":"Qi","year":"2025","journal-title":"IEEE Trans. Circuits Syst. Video Technol."},{"key":"10.1016\/j.neucom.2026.133839_bib0595","series-title":"Computer Vision \u2013 ECCV 2024","first-page":"305","article-title":"Long-term temporal context gathering for neural video compression","author":"Qi","year":"2025"},{"key":"10.1016\/j.neucom.2026.133839_bib0600","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"6111","article-title":"Motion information propagation for neural video compression","author":"Qi","year":"2023"},{"key":"10.1016\/j.neucom.2026.133839_bib0605","doi-asserted-by":"crossref","first-page":"3","DOI":"10.1016\/S0923-5965(01)00024-8","article-title":"An overview of the JPEG 2000 still image compression standard","volume":"17","author":"Rabbani","year":"2002","journal-title":"Signal Process. Image Commun."},{"key":"10.1016\/j.neucom.2026.133839_bib0610","series-title":"Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition","first-page":"4161","article-title":"Optical flow estimation using a spatial pyramid network","author":"Ranjan","year":"2017"},{"key":"10.1016\/j.neucom.2026.133839_bib0615","series-title":"Proceedings of the Asian Conference on Computer Vision (ACCV)","first-page":"3447","article-title":"Neural residual flow fields for efficient video representations","author":"Rho","year":"2022"},{"key":"10.1016\/j.neucom.2026.133839_bib0620","series-title":"Proceedings of the IEEE\/CVF International Conference on Computer Vision (ICCV)","first-page":"14479","article-title":"ELF-VC: efficient learned flexible-rate video coding","author":"Rippel","year":"2021"},{"key":"10.1016\/j.neucom.2026.133839_bib0625","doi-asserted-by":"crossref","first-page":"7311","DOI":"10.1109\/TMM.2022.3220421","article-title":"Temporal context mining for learned video compression","volume":"25","author":"Sheng","year":"2023","journal-title":"IEEE Trans. Multimed."},{"key":"10.1016\/j.neucom.2026.133839_bib0630","series-title":"IEEE Transactions on Circuits and Systems for Video Technology","first-page":"6460","article-title":"Spatial decomposition and temporal fusion based inter prediction for learned video compression","author":"Sheng","year":"2024"},{"key":"10.1016\/j.neucom.2026.133839_bib0635","doi-asserted-by":"crossref","first-page":"2285","DOI":"10.1109\/TIP.2025.3555401","article-title":"Prediction and reference quality adaptation for learned video compression","volume":"34","author":"Sheng","year":"2025","journal-title":"IEEE Trans. Image Process."},{"key":"10.1016\/j.neucom.2026.133839_bib0640","doi-asserted-by":"crossref","first-page":"5632","DOI":"10.1109\/TMM.2025.3543061","article-title":"Bi-directional deep contextual video compression","volume":"27","author":"Sheng","year":"2025","journal-title":"IEEE Trans. Multimed."},{"key":"10.1016\/j.neucom.2026.133839_bib0645","series-title":"Advances in Neural Information Processing Systems","first-page":"802","article-title":"Convolutional LSTM network: a machine learning approach for precipitation nowcasting","author":"Shi","year":"2015"},{"key":"10.1016\/j.neucom.2026.133839_bib0650","series-title":"European Conference on Computer Vision","first-page":"616","article-title":"Alphavc: high-performance and efficient learned video compression","author":"Shi","year":"2022"},{"key":"10.1016\/j.neucom.2026.133839_bib0655","series-title":"Computer Vision \u2013 ECCV 2022","first-page":"74","article-title":"Implicit neural representations for image compression","author":"Str\u00fcmpler","year":"2022"},{"key":"10.1016\/j.neucom.2026.133839_bib0660","doi-asserted-by":"crossref","first-page":"1649","DOI":"10.1109\/TCSVT.2012.2221191","article-title":"Overview of the high efficiency video coding (HEVC) standard","volume":"22","author":"Sullivan","year":"2012","journal-title":"IEEE Trans. Circuits Syst. Video Technol."},{"key":"10.1016\/j.neucom.2026.133839_bib0665","series-title":"Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (CVPR)","article-title":"PWC-net: CNNs for optical flow using pyramid, warping, and cost volume","author":"Sun","year":"2018"},{"key":"10.1016\/j.neucom.2026.133839_bib0670","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"12553","article-title":"Neural video compression with context modulation","author":"Tang","year":"2025"},{"key":"10.1016\/j.neucom.2026.133839_bib0675","series-title":"Proceedings of the AAAI Conference on Artificial Intelligence","first-page":"5118","article-title":"Offline and online optical flow enhancement for deep video compression","author":"Tang","year":"2024"},{"key":"10.1016\/j.neucom.2026.133839_bib0680","series-title":"Image Processing Algorithms and Techniques","first-page":"220","article-title":"Overview of the JPEG (ISO\/CCITT) still image compression standard","author":"Wallace","year":"1990"},{"key":"10.1016\/j.neucom.2026.133839_bib0685","series-title":"ICASSP 2025 - 2025 IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP)","first-page":"1","article-title":"M3-CVC: controllable video compression with multimodal generative models","author":"Wan","year":"2025"},{"key":"10.1016\/j.neucom.2026.133839_bib0690","doi-asserted-by":"crossref","first-page":"780","DOI":"10.1109\/TIP.2024.3349859","article-title":"Exploring long- and short-range temporal information for learned video compression","volume":"33","author":"Wang","year":"2024","journal-title":"IEEE Trans. Image Process."},{"key":"10.1016\/j.neucom.2026.133839_bib0695","doi-asserted-by":"crossref","first-page":"1855","DOI":"10.1109\/TMM.2023.3289763","article-title":"Learned video compression via heterogeneous deformable compensation network","volume":"26","author":"Wang","year":"2024","journal-title":"IEEE Trans. Multimed."},{"key":"10.1016\/j.neucom.2026.133839_bib0700","series-title":"2016 IEEE International Conference on Image Processing (ICIP)","first-page":"1509","article-title":"MCL-JCV: a JND-based H.264\/AVC video quality assessment dataset","author":"Wang","year":"2016"},{"key":"10.1016\/j.neucom.2026.133839_bib0705","series-title":"2024 Picture Coding Symposium (PCS)","first-page":"1","article-title":"Temporal enhanced hybrid neural representation for video compression","author":"Wang","year":"2024"},{"key":"10.1016\/j.neucom.2026.133839_bib0710","series-title":"Image and Graphics","first-page":"360","article-title":"Learning to fuse residual and conditional information for video compression and reconstruction","author":"Wang","year":"2023"},{"key":"10.1016\/j.neucom.2026.133839_bib0715","series-title":"Proceedings of the IEEE\/CVF International Conference on Computer Vision (ICCV) Workshops","first-page":"1905","article-title":"Real-ESRGAN: training real-world blind super-resolution with pure synthetic data","author":"Wang","year":"2021"},{"key":"10.1016\/j.neucom.2026.133839_bib0720","series-title":"Proceedings of the European Conference on Computer Vision (ECCV) Workshops","article-title":"ESRGAN: enhanced super-resolution generative adversarial networks","author":"Wang","year":"2018"},{"key":"10.1016\/j.neucom.2026.133839_bib0725","series-title":"ICASSP 2024 - 2024 IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP)","first-page":"3715","article-title":"Learned video compression with spatial-temporal optimization","author":"Wang","year":"2024"},{"key":"10.1016\/j.neucom.2026.133839_bib0730","series-title":"2023 IEEE International Conference on Image Processing (ICIP)","first-page":"3175","article-title":"FGC-VC: flow-guided context video compression","author":"Wang","year":"2023"},{"key":"10.1016\/j.neucom.2026.133839_bib0735","doi-asserted-by":"crossref","DOI":"10.1016\/j.eswa.2024.123322","article-title":"Temporal context video compression with flow-guided feature prediction","volume":"247","author":"Wang","year":"2024","journal-title":"Expert Syst. Appl."},{"key":"10.1016\/j.neucom.2026.133839_bib0740","series-title":"Advances in Neural Information Processing Systems","article-title":"PredRNN: recurrent neural networks for predictive learning using spatiotemporal LSTMs","author":"Wang","year":"2017"},{"key":"10.1016\/j.neucom.2026.133839_bib0745","doi-asserted-by":"crossref","first-page":"560","DOI":"10.1109\/TCSVT.2003.815165","article-title":"Overview of the H.264\/AVC video coding standard","volume":"13","author":"Wiegand","year":"2003","journal-title":"IEEE Trans. Circuits Syst. Video Technol."},{"key":"10.1016\/j.neucom.2026.133839_bib0750","series-title":"Association for Computing Machinery","first-page":"2584","article-title":"QS-NeRV: real-time quality-scalable decoding with neural representation for videos","author":"Wu","year":"2024"},{"key":"10.1016\/j.neucom.2026.133839_bib0755","series-title":"Proceedings of the European Conference on Computer Vision (ECCV)","article-title":"Video compression through image interpolation","author":"Wu","year":"2018"},{"key":"10.1016\/j.neucom.2026.133839_bib0760","doi-asserted-by":"crossref","DOI":"10.1016\/j.cviu.2024.104127","article-title":"Deep video compression based on long-range temporal context learning","volume":"248","author":"Wu","year":"2024","journal-title":"Comput. Vis. Image Underst."},{"key":"10.1016\/j.neucom.2026.133839_bib0765","doi-asserted-by":"crossref","first-page":"4386","DOI":"10.1109\/TMM.2025.3535367","article-title":"End-to-end deep video compression based on hierarchical temporal context learning","volume":"27","author":"Wu","year":"2025","journal-title":"IEEE Trans. Multimed."},{"key":"10.1016\/j.neucom.2026.133839_bib0770","series-title":"2023 IEEE International Conference on Visual Communications and Image Processing (VCIP)","first-page":"1","article-title":"Rate adaptation for learned two-layer B-frame coding without signaling motion information","author":"Xie","year":"2023"},{"key":"10.1016\/j.neucom.2026.133839_bib0775","series-title":"Proceedings of the 29th ACM International Conference on Multimedia","first-page":"162","article-title":"Enhanced invertible encoding for learned image compression","author":"Xie","year":"2021"},{"key":"10.1016\/j.neucom.2026.133839_bib0780","doi-asserted-by":"crossref","DOI":"10.1016\/j.patcog.2024.110465","article-title":"IBVC: interpolation-driven B-frame video compression","volume":"153","author":"Xu","year":"2024","journal-title":"Pattern Recognit."},{"key":"10.1016\/j.neucom.2026.133839_bib0785","first-page":"1","article-title":"Deep video coding with bit-depth scalability","author":"Xu","year":"2026","journal-title":"IEEE Trans. Multimed."},{"key":"10.1016\/j.neucom.2026.133839_bib0790","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"23019","article-title":"DS-NeRV: implicit neural video representation with decomposed static and dynamic codes","author":"Yan","year":"2024"},{"key":"10.1016\/j.neucom.2026.133839_bib0795","doi-asserted-by":"crossref","first-page":"2125","DOI":"10.1109\/LSP.2024.3443516","article-title":"Multi-scale motion alignment and frame reconstruction for efficient deep video compression","volume":"31","author":"Yang","year":"2024","journal-title":"IEEE Signal Process. Lett."},{"key":"10.1016\/j.neucom.2026.133839_bib0800","series-title":"2024 Data Compression Conference (DCC)","first-page":"382","article-title":"UCVC: a unified contextual video compression framework with joint P-frame and B-frame coding","author":"Yang","year":"2024"},{"key":"10.1016\/j.neucom.2026.133839_bib0805","doi-asserted-by":"crossref","DOI":"10.1145\/3661824","article-title":"Learned video compression with adaptive temporal prior and decoded motion-aided quality enhancement","volume":"20","author":"Yang","year":"2024","journal-title":"ACM Trans. Multimedia Comput. Commun. Appl."},{"key":"10.1016\/j.neucom.2026.133839_bib0810","series-title":"ICASSP 2024 - 2024 IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP)","first-page":"2860","article-title":"Improving learned video compression by exploring spatial redundancy","author":"Yang","year":"2024"},{"key":"10.1016\/j.neucom.2026.133839_bib0815","doi-asserted-by":"crossref","DOI":"10.1145\/3703914","article-title":"Adaptive prediction structure for learned video compression","volume":"21","author":"Yang","year":"2025","journal-title":"ACM Trans. Multimedia Comput. Commun. Appl."},{"key":"10.1016\/j.neucom.2026.133839_bib0820","series-title":"2024 IEEE International Conference on Image Processing (ICIP)","first-page":"3723","article-title":"Learning-based video compression with continuously variable bitrate coding","author":"Yang","year":"2024"},{"key":"10.1016\/j.neucom.2026.133839_bib0825","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition","first-page":"6628","article-title":"Learning for video compression with hierarchical quality and recurrent enhancement","author":"Yang","year":"2020"},{"key":"10.1016\/j.neucom.2026.133839_bib0830","doi-asserted-by":"crossref","first-page":"388","DOI":"10.1109\/JSTSP.2020.3043590","article-title":"Learning for video compression with recurrent auto-encoder and recurrent probability model","volume":"15","author":"Yang","year":"2021","journal-title":"IEEE J. Sel. Top. Signal Process."},{"key":"10.1016\/j.neucom.2026.133839_bib0835","series-title":"2019 IEEE International Conference on Multimedia and Expo (ICME)","first-page":"532","article-title":"Quality-gated convolutional LSTM for enhancing compressed video","author":"Yang","year":"2019"},{"key":"10.1016\/j.neucom.2026.133839_bib0840","series-title":"IJCAI","first-page":"1537","article-title":"Perceptual learned video compression with recurrent conditional GAN","author":"Yang","year":"2022"},{"key":"10.1016\/j.neucom.2026.133839_bib0845","doi-asserted-by":"crossref","first-page":"2410","DOI":"10.1109\/TCSVT.2022.3222418","article-title":"Advancing learned video compression with in-loop frame prediction","volume":"33","author":"Yang","year":"2023","journal-title":"IEEE Trans. Circuits Syst. Video Technol."},{"key":"10.1016\/j.neucom.2026.133839_bib0850","doi-asserted-by":"crossref","first-page":"9908","DOI":"10.1109\/TPAMI.2023.3260684","article-title":"Insights from generative modeling for neural video compression","volume":"45","author":"Yang","year":"2023","journal-title":"IEEE Trans. Pattern Anal. Mach. Intell."},{"key":"10.1016\/j.neucom.2026.133839_bib0855","series-title":"Proc. ICML Workshop Invertible Neural Netw. Normalizing Flows, Explicit Likelihood Models","article-title":"Deep generative video compression with temporal autoregressive transforms","author":"Yang","year":"2020"},{"key":"10.1016\/j.neucom.2026.133839_bib0860","series-title":"Proceedings of the 32nd ACM International Conference on Multimedia","first-page":"11244","article-title":"Deep video compression with scaled hierarchical bi-directional motion model","author":"Ye","year":"2024"},{"key":"10.1016\/j.neucom.2026.133839_bib0865","doi-asserted-by":"crossref","first-page":"808","DOI":"10.1109\/TBC.2025.3587532","article-title":"Multi-domain spatial-temporal redundancy mining for efficient learned video compression","volume":"71","author":"Yuan","year":"2025","journal-title":"IEEE Trans. Broadcast."},{"key":"10.1016\/j.neucom.2026.133839_bib0870","doi-asserted-by":"crossref","first-page":"974","DOI":"10.1109\/TIP.2021.3138300","article-title":"End-to-end rate-distortion optimized learned hierarchical bi-directional video compression","volume":"31","author":"Y\u0131lmaz","year":"2022","journal-title":"IEEE Trans. Image Process."},{"key":"10.1016\/j.neucom.2026.133839_bib0875","series-title":"2023 IEEE International Conference on Image Processing (ICIP)","first-page":"2475","article-title":"Multi-scale deformable alignment and content-adaptive inference for flexible-rate bi-directional video compression","author":"Y\u0131lmaz","year":"2023"},{"key":"10.1016\/j.neucom.2026.133839_bib0880","series-title":"Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (CVPR)","article-title":"The unreasonable effectiveness of deep features as a perceptual metric","author":"Zhang","year":"2018"},{"key":"10.1016\/j.neucom.2026.133839_bib0885","series-title":"2021 International Conference on Visual Communications and Image Processing (VCIP)","first-page":"1","article-title":"DVC-P: deep video compression with perceptual optimizations","author":"Zhang","year":"2021"},{"key":"10.1016\/j.neucom.2026.133839_bib0890","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"2556","article-title":"Boosting neural representations for videos with a conditional decoder","author":"Zhang","year":"2024"},{"key":"10.1016\/j.neucom.2026.133839_bib0895","series-title":"2024 IEEE International Conference on Image Processing (ICIP)","first-page":"3681","article-title":"Optimized decoupled structure with non-local attention for deep image compression","author":"Zhang","year":"2024"},{"key":"10.1016\/j.neucom.2026.133839_bib0900","doi-asserted-by":"crossref","first-page":"3067","DOI":"10.1109\/TCSVT.2023.3313974","article-title":"End-to-end learning-based image compression with a decoupled framework","volume":"34","author":"Zhang","year":"2024","journal-title":"IEEE Trans. Circuits Syst. Video Technol."},{"key":"10.1016\/j.neucom.2026.133839_bib0905","author":"Zhao"},{"key":"10.1016\/j.neucom.2026.133839_bib0910","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"2031","article-title":"DNeRV: modeling inherent dynamics via difference neural representation for videos","author":"Zhao","year":"2023"},{"key":"10.1016\/j.neucom.2026.133839_bib0915","series-title":"Proceedings of the IEEE\/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","first-page":"19103","article-title":"PNeRV: enhancing spatial consistency via pyramidal neural representation for videos","author":"Zhao","year":"2024"},{"key":"10.1016\/j.neucom.2026.133839_bib0920","series-title":"Proceedings of the 30th ACM International Conference on Multimedia","first-page":"3045","article-title":"Learning-based video coding with joint deep compression and enhancement","author":"Zhao","year":"2022"},{"key":"10.1016\/j.neucom.2026.133839_bib0925","series-title":"2024 International Joint Conference on Neural Networks (IJCNN)","first-page":"1","article-title":"PFR-VC: learning-based video compression framework with predicted frame refinement","author":"Zhou","year":"2024"},{"key":"10.1016\/j.neucom.2026.133839_bib0930","series-title":"2022 IEEE International Conference on Image Processing (ICIP)","first-page":"1206","article-title":"Flexible-rate learned hierarchical bi-directional video compression with motion refinement and frame-level bit allocation","author":"\u00c7etin","year":"2022"}],"container-title":["Neurocomputing"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/api.elsevier.com\/content\/article\/PII:S0925231226012361?httpAccept=text\/xml","content-type":"text\/xml","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/api.elsevier.com\/content\/article\/PII:S0925231226012361?httpAccept=text\/plain","content-type":"text\/plain","content-version":"vor","intended-application":"text-mining"}],"deposited":{"date-parts":[[2026,7,13]],"date-time":"2026-07-13T11:51:08Z","timestamp":1783943468000},"score":1,"resource":{"primary":{"URL":"https:\/\/linkinghub.elsevier.com\/retrieve\/pii\/S0925231226012361"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2026,9]]},"references-count":184,"alternative-id":["S0925231226012361"],"URL":"https:\/\/doi.org\/10.1016\/j.neucom.2026.133839","relation":{},"ISSN":["0925-2312"],"issn-type":[{"value":"0925-2312","type":"print"}],"subject":[],"published":{"date-parts":[[2026,9]]},"assertion":[{"value":"Elsevier","name":"publisher","label":"This article is maintained by"},{"value":"End-to-end learned video compression: A comprehensive review","name":"articletitle","label":"Article Title"},{"value":"Neurocomputing","name":"journaltitle","label":"Journal Title"},{"value":"https:\/\/doi.org\/10.1016\/j.neucom.2026.133839","name":"articlelink","label":"CrossRef DOI link to publisher maintained version"},{"value":"article","name":"content_type","label":"Content Type"},{"value":"\u00a9 2026 Published by Elsevier B.V.","name":"copyright","label":"Copyright"}],"article-number":"133839"}}