{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2025,7,3]],"date-time":"2025-07-03T13:34:57Z","timestamp":1751549697008,"version":"3.37.3"},"reference-count":50,"publisher":"Springer Science and Business Media LLC","issue":"6","license":[{"start":{"date-parts":[[2021,1,24]],"date-time":"2021-01-24T00:00:00Z","timestamp":1611446400000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springer.com\/tdm"},{"start":{"date-parts":[[2021,1,24]],"date-time":"2021-01-24T00:00:00Z","timestamp":1611446400000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springer.com\/tdm"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Int. J. Mach. Learn. &amp; Cyber."],"published-print":{"date-parts":[[2021,6]]},"DOI":"10.1007\/s13042-020-01272-7","type":"journal-article","created":{"date-parts":[[2021,1,24]],"date-time":"2021-01-24T09:04:38Z","timestamp":1611479078000},"page":"1753-1768","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":2,"title":["Sample-based online learning for bi-regular hinge loss"],"prefix":"10.1007","volume":"12","author":[{"ORCID":"https:\/\/orcid.org\/0000-0002-0082-462X","authenticated-orcid":false,"given":"Wei","family":"Xue","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Ping","family":"Zhong","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Wensheng","family":"Zhang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Gaohang","family":"Yu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Yebin","family":"Chen","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2021,1,24]]},"reference":[{"issue":"4","key":"1272_CR1","doi-asserted-by":"publisher","first-page":"1746","DOI":"10.1109\/TAC.2018.2860546","volume":"64","author":"M Akbari","year":"2019","unstructured":"Akbari M, Gharesifard B, Linder T (2019) Individual regret bounds for the distributed online alternating direction method of multipliers. IEEE Trans Autom Control 64(4):1746\u20131752","journal-title":"IEEE Trans Autom Control"},{"key":"1272_CR2","first-page":"319","volume":"2","author":"D Angluin","year":"1988","unstructured":"Angluin D (1988) Queries and concept learning. Mach Learn 2:319\u2013342","journal-title":"Mach Learn"},{"key":"1272_CR3","doi-asserted-by":"publisher","first-page":"141","DOI":"10.1093\/imanum\/8.1.141","volume":"8","author":"J Barzilai","year":"1988","unstructured":"Barzilai J, Borwein JM (1988) Two-point step size gradient methods. IMA J Numer Anal 8:141\u2013148","journal-title":"IMA J Numer Anal"},{"issue":"1","key":"1272_CR4","doi-asserted-by":"publisher","first-page":"1","DOI":"10.1561\/2200000016","volume":"3","author":"S Boyd","year":"2010","unstructured":"Boyd S, Parikh N, Chu E, Peleato B, Eckstein J (2010) Distributed optimization and statistical learning via the alternating direction method of multipliers. Found Trends Mach Learn 3(1):1\u2013122","journal-title":"Found Trends Mach Learn"},{"key":"1272_CR5","doi-asserted-by":"publisher","DOI":"10.1007\/978-3-642-20192-9","volume-title":"Statistics for High-dimensional Data: Methods, Theory and Applications","author":"P Buhlmann","year":"2011","unstructured":"Buhlmann P, van de Geer S (2011) Statistics for High-dimensional Data: Methods, Theory and Applications. Springer, Berlin"},{"issue":"5","key":"1272_CR6","doi-asserted-by":"publisher","first-page":"877","DOI":"10.1007\/s00041-008-9045-x","volume":"14","author":"EJ Cand\u00e8s","year":"2008","unstructured":"Cand\u00e8s EJ, Wakin MB, Boyd SP (2008) Enhancing sparsity by reweighted $$l_1$$ minimization. J Fourier Anal Appl 14(5):877\u2013905","journal-title":"J Fourier Anal Appl"},{"key":"1272_CR7","first-page":"1369","volume":"9","author":"KW Chang","year":"2008","unstructured":"Chang KW, Hsieh CJ, Lin CJ (2008) Coordinate descent method for large-scale l2-loss linear support vector machines. J Mach Learn Res 9:1369\u20131398","journal-title":"J Mach Learn Res"},{"issue":"3","key":"1272_CR8","doi-asserted-by":"publisher","first-page":"27:1","DOI":"10.1145\/1961189.1961199","volume":"2","author":"CC Chang","year":"2011","unstructured":"Chang CC, Lin CJ (2011) LIBSVM: a library for support vector machines. ACM Trans Intell Syst Technol 2(3):27:1\u201327:27","journal-title":"ACM Trans Intell Syst Technol"},{"key":"1272_CR9","doi-asserted-by":"publisher","first-page":"803","DOI":"10.1007\/s10462-018-9614-6","volume":"52","author":"VK Chauhan","year":"2019","unstructured":"Chauhan VK, Dahiya K, Sharma A (2019) Problem formulations and solvers in linear SVM: a review. Artif Intell Rev 52:803\u2013855","journal-title":"Artif Intell Rev"},{"key":"1272_CR10","doi-asserted-by":"publisher","first-page":"1541","DOI":"10.1007\/s13042-019-01055-9","volume":"11","author":"VK Chauhan","year":"2020","unstructured":"Chauhan VK, Sharma A, Dahiya K (2020) Stochastic trust region inexact Newton method for large-scale machine learning. Int J Mach Learn and Cybern 11:1541\u20131555","journal-title":"Int J Mach Learn and Cybern"},{"issue":"11","key":"1272_CR11","doi-asserted-by":"publisher","first-page":"5974","DOI":"10.1109\/TAC.2017.2705559","volume":"62","author":"K Cohen","year":"2017","unstructured":"Cohen K, Nedi\u0107 A, Srikant R (2017) On projected stochastic gradient descent algorithm with weighted averaging for least squares regression. IEEE Trans Autom Control 62(11):5974\u20135981","journal-title":"IEEE Trans Autom Control"},{"issue":"3","key":"1272_CR12","first-page":"273","volume":"20","author":"C Cortes","year":"1995","unstructured":"Cortes C, Vapnik V (1995) Support-vector networks. Mach Learn 20(3):273\u2013297","journal-title":"Mach Learn"},{"issue":"2","key":"1272_CR13","doi-asserted-by":"publisher","first-page":"201","DOI":"10.1016\/j.jco.2009.01.002","volume":"25","author":"C De Mol","year":"2009","unstructured":"De Mol C, De Vito E, Rosasco L (2009) Elastic-net regularization in learning theory. J Complex 25(2):201\u2013230","journal-title":"J Complex"},{"key":"1272_CR14","unstructured":"Duchi JC, Shalev-Shwartz S, Singer Y, Tewari A (2010) Composite Objective Mirror Descent. In: Proceedings of the 23rd annual conference on learning theory, pp 14\u201326"},{"issue":"1","key":"1272_CR15","first-page":"17","volume":"2","author":"D Gabay","year":"1976","unstructured":"Gabay D, Mercier B (1976) A dual algorithm for the solution of nonlinear variational problems via finite element approximation. Comput Optim Appl 2(1):17\u201340","journal-title":"Comput Optim Appl"},{"key":"1272_CR16","volume-title":"Computers and intractability: a guide to the theory of NP-completeness","author":"MR Garey","year":"1979","unstructured":"Garey MR, Johnson DS (1979) Computers and intractability: a guide to the theory of NP-completeness. Freeman, New York"},{"key":"1272_CR17","first-page":"41","volume":"9","author":"R Glowinski","year":"1975","unstructured":"Glowinski R, Marrocco A (1975) Sur l\u2019approximation, par\u00e9l\u00e9ments finis d\u2019ordre un, et lan r\u00e9solution, par p\u00e9nalisation-dualit\u00e9, d\u2019une classe de probl\u00e9mes de Dirichlet non lin\u00e9aires. Rev Fr Automat Infor 9:41\u201376","journal-title":"Rev Fr Automat Infor"},{"key":"1272_CR18","volume-title":"Machine learning for multimedia content analysis","author":"Y Gong","year":"2007","unstructured":"Gong Y, Xu W (2007) Machine learning for multimedia content analysis. Springer Science & Business Media, New York"},{"key":"1272_CR19","doi-asserted-by":"crossref","unstructured":"Hajewski J, Oliveira S, Stewart D (2018) Smoothed hinge loss and l1 support vector machines. In: Proceedings of the 2018 IEEE international conference on data mining workshops, pp 1217\u20131223","DOI":"10.1109\/ICDMW.2018.00174"},{"issue":"2","key":"1272_CR20","doi-asserted-by":"publisher","first-page":"700","DOI":"10.1137\/110836936","volume":"50","author":"B He","year":"2012","unstructured":"He B, Yuan X (2012) On the $$O(1\/n)$$ convergence rate of the Douglas-Rachford alternating direction method. SIAM J Numer Anal 50(2):700\u2013709","journal-title":"SIAM J Numer Anal"},{"key":"1272_CR21","doi-asserted-by":"crossref","unstructured":"Hsieh CJ, Chang KW, Lin CJ, Keerthi SS, Sundararajan S (2008) A dual coordinate descent method for large-scale linear SVM. In: Proceedings of the 36th international conference on machine learning, pp 408\u2013415","DOI":"10.1145\/1390156.1390208"},{"key":"1272_CR22","doi-asserted-by":"crossref","unstructured":"Huang F, Chen S, Huang H (2019) Faster stochastic alternating direction method of multipliers for nonconvex optimization. In: Proceedings of the 36th international conference on machine learning, pp 2839\u20132848","DOI":"10.24963\/ijcai.2019\/354"},{"key":"1272_CR23","doi-asserted-by":"crossref","unstructured":"Joachims T (2006) Training linear SVMs in linear time. In: Proceedings of the 12th ACM SIGKDD international conference on knowledge discovery and data mining, pp 217\u2013226","DOI":"10.1145\/1150402.1150429"},{"key":"1272_CR24","unstructured":"Johnson R, Zhang T (2013) Accelerating stochastic gradient descent using predictive variance reduction. In: Advances in Neural Information Processing Systems, pp 315-323"},{"key":"1272_CR25","doi-asserted-by":"publisher","first-page":"179575","DOI":"10.1109\/ACCESS.2019.2954859","volume":"7","author":"ZA Khan","year":"2019","unstructured":"Khan ZA, Zubair S, Alquhayz H, Azeem M, Ditta A (2019) Design of momentum fractional stochastic gradient descent for recommender systems. IEEE Access 7:179575\u2013179590","journal-title":"IEEE Access"},{"key":"1272_CR26","first-page":"627","volume":"9","author":"CJ Lin","year":"2008","unstructured":"Lin CJ, Weng RC, Sathiya Keerthi S (2008) Trust region Newton method for large-scale logistic regression. J Mach Learn Res 9:627\u2013650","journal-title":"J Mach Learn Res"},{"key":"1272_CR27","doi-asserted-by":"crossref","unstructured":"Liu Y, Shang F, Cheng J (2017) Accelerated variance reduced stochastic ADMM. In: Proceedings of the 31st AAAI conference on artificial intelligence, pp 2287-2293","DOI":"10.1609\/aaai.v31i1.10843"},{"issue":"4","key":"1272_CR28","first-page":"285","volume":"2","author":"N Littlestone","year":"1988","unstructured":"Littlestone N (1988) Learning quickly when irrelevant attributes abound: a new linear-threshold algorithm. Mach Learn 2(4):285\u2013318","journal-title":"Mach Learn"},{"issue":"2","key":"1272_CR29","doi-asserted-by":"publisher","first-page":"857","DOI":"10.1007\/s10462-017-9611-1","volume":"52","author":"J Nalepa","year":"2019","unstructured":"Nalepa J, Kawulok M (2019) Selecting training sets for support vector machines: a review. Artif Intell Rev 52(2):857\u2013900","journal-title":"Artif Intell Rev"},{"key":"1272_CR30","volume-title":"Encyclopedia of Machine Learning","author":"C Sammut","year":"2011","unstructured":"Sammut C, Webb GI (2011) Encyclopedia of Machine Learning. Springer Science & Business Media, New York"},{"issue":"2","key":"1272_CR31","doi-asserted-by":"publisher","first-page":"107","DOI":"10.1561\/2200000018","volume":"4","author":"S Shalev-Shwartz","year":"2012","unstructured":"Shalev-Shwartz S (2012) Online learning and online convex optimization. Found Trends Mach Learn 4(2):107\u2013194","journal-title":"Found Trends Mach Learn"},{"issue":"1","key":"1272_CR32","doi-asserted-by":"publisher","first-page":"3","DOI":"10.1007\/s10107-010-0420-4","volume":"127","author":"S Shalev-Shwartz","year":"2011","unstructured":"Shalev-Shwartz S, Singer Y, Srebro N, Cotter A (2011) Pegasos: primal estimated sub-gradient solver for SVM. Math Program 127(1):3\u201330","journal-title":"Math Program"},{"key":"1272_CR33","doi-asserted-by":"publisher","first-page":"11173","DOI":"10.1007\/s00521-019-04627-6","volume":"32","author":"M Singla","year":"2020","unstructured":"Singla M, Shukla KK (2020) Robust statistics-based support vector machine and its variants: a survey. Neural Comput Appl 32:11173\u201311194","journal-title":"Neural Comput Appl"},{"key":"1272_CR34","doi-asserted-by":"publisher","first-page":"64533","DOI":"10.1109\/ACCESS.2019.2915970","volume":"7","author":"T Song","year":"2019","unstructured":"Song T, Li D, Liu Z, Yang W (2019) Online ADMM-based extreme learning machine for sparse supervised learning. IEEE Access 7:64533\u201364544","journal-title":"IEEE Access"},{"key":"1272_CR35","unstructured":"Suzuki T (2013) Dual averaging and proximal gradient descent for online alternating direction multiplier method. In: Proceedings of the 30th international conference on machine learning, pp 392\u2013400"},{"key":"1272_CR36","unstructured":"Tan C, Ma S, Dai Y H, Qian Y (2016) Barzilai-borwein step size for stochastic gradient descent. Advances in neural information processing systems, pp 685\u2013693"},{"key":"1272_CR37","doi-asserted-by":"publisher","DOI":"10.1007\/978-1-4757-2440-0","volume-title":"The nature of statistical learning theory","author":"V Vapnik","year":"1995","unstructured":"Vapnik V (1995) The nature of statistical learning theory. Springer, New York"},{"key":"1272_CR38","first-page":"589","volume":"16","author":"L Wang","year":"2006","unstructured":"Wang L, Zhu J, Zou H (2006) The doubly regularized support vector machine. Stat Sinica 16:589\u2013615","journal-title":"Stat Sinica"},{"issue":"3","key":"1272_CR39","doi-asserted-by":"publisher","first-page":"412","DOI":"10.1093\/bioinformatics\/btm579","volume":"2","author":"L Wang","year":"2008","unstructured":"Wang L, Zhu J, Zou H (2008) Hybrid huberized support vector machines for microarray classification and gene selection. Bioinformatics 2(3):412\u2013419","journal-title":"Bioinformatics"},{"issue":"5","key":"1272_CR40","doi-asserted-by":"publisher","first-page":"802","DOI":"10.1109\/TCSVT.2013.2290574","volume":"24","author":"Z Wang","year":"2014","unstructured":"Wang Z, Hu R, Wang S, Jiang J (2014) Face hallucination via weighted adaptive sparse regularization. IEEE Trans Circuits Syst Video Technol 24(5):802\u2013813","journal-title":"IEEE Trans Circuits Syst Video Technol"},{"key":"1272_CR41","unstructured":"Xiao L (2009) Dual averaging methods for regularized stochastic learning and online optimization. In: Advances in neural information processing systems, pp 2116\u20132124"},{"issue":"6","key":"1272_CR42","doi-asserted-by":"publisher","first-page":"1529","DOI":"10.1007\/s13042-018-0832-7","volume":"10","author":"Z Xie","year":"2019","unstructured":"Xie Z, Li Y (2019) Large-scale support vector regression with budgeted stochastic gradient descent. Int J Mach Learn Cybern 10(6):1529\u20131541","journal-title":"Int J Mach Learn Cybern"},{"issue":"4","key":"1272_CR43","doi-asserted-by":"publisher","first-page":"989","DOI":"10.1007\/s10044-015-0485-z","volume":"19","author":"Y Xu","year":"2016","unstructured":"Xu Y, Akrotirianakis I, Chakraborty A (2016) Proximal gradient method for huberized support vector machine. Pattern Anal Appl 19(4):989\u20131005","journal-title":"Pattern Anal Appl"},{"issue":"2","key":"1272_CR44","doi-asserted-by":"publisher","first-page":"438","DOI":"10.1109\/TNNLS.2016.2514413","volume":"28","author":"W Xue","year":"2017","unstructured":"Xue W, Zhang W (2017) Learning a coupled linearized method in online setting. IEEE Trans Neural Netw Learn Syst 28(2):438\u2013450","journal-title":"IEEE Trans Neural Netw Learn Syst"},{"key":"1272_CR45","doi-asserted-by":"publisher","first-page":"420","DOI":"10.1016\/j.neucom.2017.04.044","volume":"260","author":"E Zamora","year":"2017","unstructured":"Zamora E, Sossa H (2017) Dendrite morphological neurons trained by stochastic gradient descent. Neurocomputing 260:420\u2013431","journal-title":"Neurocomputing"},{"key":"1272_CR46","unstructured":"Zhao P, Zhang T (2015) Stochastic optimization with importance sampling for regularized loss minimization. In: Proceedings of the 20th international conference on machine learning"},{"key":"1272_CR47","unstructured":"Zhu J, Rosset S, Hastie T, Tibshirani R (2004) 1-norm support vector machines. In: Advances in neural information processing systems, pp 49\u201356"},{"key":"1272_CR48","unstructured":"Zinkevich M (2003) Online convex programming and generalized infinitesimal gradient ascent. In: Proceedings of the 20th international conference on machine learning, pp 928\u2013936"},{"issue":"2","key":"1272_CR49","doi-asserted-by":"publisher","first-page":"301","DOI":"10.1111\/j.1467-9868.2005.00503.x","volume":"67","author":"H Zou","year":"2005","unstructured":"Zou H, Hastie T (2005) Regularization and variable selection via the elastic net. J R Stat Soc B 67(2):301\u2013320","journal-title":"J R Stat Soc B"},{"issue":"4","key":"1272_CR50","doi-asserted-by":"publisher","first-page":"1733","DOI":"10.1214\/08-AOS625","volume":"37","author":"H Zou","year":"2009","unstructured":"Zou H, Zhang H (2009) On the adaptive elastic-net with a diverging number of parameters. Ann Stat 37(4):1733\u20131751","journal-title":"Ann Stat"}],"container-title":["International Journal of Machine Learning and Cybernetics"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s13042-020-01272-7.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s13042-020-01272-7\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s13042-020-01272-7.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2022,12,13]],"date-time":"2022-12-13T02:29:22Z","timestamp":1670898562000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s13042-020-01272-7"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2021,1,24]]},"references-count":50,"journal-issue":{"issue":"6","published-print":{"date-parts":[[2021,6]]}},"alternative-id":["1272"],"URL":"https:\/\/doi.org\/10.1007\/s13042-020-01272-7","relation":{},"ISSN":["1868-8071","1868-808X"],"issn-type":[{"type":"print","value":"1868-8071"},{"type":"electronic","value":"1868-808X"}],"subject":[],"published":{"date-parts":[[2021,1,24]]},"assertion":[{"value":"10 January 2020","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"29 December 2020","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"24 January 2021","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}}]}}