{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,4,14]],"date-time":"2026-04-14T02:05:38Z","timestamp":1776132338395,"version":"3.50.1"},"reference-count":28,"publisher":"Springer Science and Business Media LLC","issue":"6","license":[{"start":{"date-parts":[[2020,11,13]],"date-time":"2020-11-13T00:00:00Z","timestamp":1605225600000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springer.com\/tdm"},{"start":{"date-parts":[[2020,11,13]],"date-time":"2020-11-13T00:00:00Z","timestamp":1605225600000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springer.com\/tdm"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Appl Intell"],"published-print":{"date-parts":[[2021,6]]},"DOI":"10.1007\/s10489-020-01906-x","type":"journal-article","created":{"date-parts":[[2020,11,13]],"date-time":"2020-11-13T00:06:22Z","timestamp":1605225982000},"page":"3405-3420","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":11,"title":["An adaptive adjustment strategy for bolt posture errors based on an improved reinforcement learning algorithm"],"prefix":"10.1007","volume":"51","author":[{"ORCID":"https:\/\/orcid.org\/0000-0001-6579-2190","authenticated-orcid":false,"given":"Wentao","family":"Luo","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Jianfu","family":"Zhang","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Pingfa","family":"Feng","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Haochen","family":"Liu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Dingwen","family":"Yu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Zhijun","family":"Wu","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2020,11,13]]},"reference":[{"key":"1906_CR1","doi-asserted-by":"publisher","first-page":"81","DOI":"10.1016\/j.ins.2017.04.012","volume":"405","author":"UR Acharya","year":"2017","unstructured":"Acharya UR, Fujita H, Lih OS, Hagiwara Y, Tan JH, Adam M (2017) Automated detection of arrhythmias using different intervals of tachycardia ECG segments with convolutional neural network. Inf Sci 405:81\u201390. https:\/\/doi.org\/10.1016\/j.ins.2017.04.012","journal-title":"Inf Sci"},{"key":"1906_CR2","doi-asserted-by":"publisher","first-page":"97","DOI":"10.1016\/j.compbiomed.2015.12.006","volume":"69","author":"VK Sudarshan","year":"2016","unstructured":"Sudarshan VK, Mookiah MRK, Acharya UR, Chandran V, Molinari F, Fujita H, Ng KH (2016) Application of wavelet techniques for cancer diagnosis using ultrasound images: a review. Comput Biol Med 69:97\u2013111. https:\/\/doi.org\/10.1016\/j.compbiomed.2015.12.006","journal-title":"Comput Biol Med"},{"key":"1906_CR3","doi-asserted-by":"publisher","first-page":"190","DOI":"10.1016\/j.ins.2017.06.027","volume":"415-416","author":"UR Acharya","year":"2017","unstructured":"Acharya UR, Fujita H, Oh SL, Hagiwara Y, Tan JH, Adam M (2017) Application of deep convolutional neural network for automated detection of myocardial infarction using ECG signals. Inf Sci 415-416:190\u2013198. https:\/\/doi.org\/10.1016\/j.ins.2017.06.027","journal-title":"Inf Sci"},{"issue":"3","key":"1906_CR4","doi-asserted-by":"publisher","first-page":"1704","DOI":"10.1109\/TFUZZ.2017.2744605","volume":"26","author":"N Capuano","year":"2018","unstructured":"Capuano N, Chiclana F, Fujita H, Herrera-Viedma E, Loia V (2018) Fuzzy group decision making with incomplete information guided by social influence. IEEE Trans Fuzzy Syst 26(3):1704\u20131718. https:\/\/doi.org\/10.1109\/TFUZZ.2017.2744605","journal-title":"IEEE Trans Fuzzy Syst"},{"issue":"7","key":"1906_CR5","doi-asserted-by":"publisher","first-page":"2793","DOI":"10.1007\/s10489-018-01396-y","volume":"49","author":"E Protopapadakis","year":"2019","unstructured":"Protopapadakis E, Voulodimos A, Doulamis A, Doulamis N, Stathaki T (2019) Automatic crack detection for tunnel inspection using deep learning and heuristic image post-processing. Appl Intell 49(7):2793\u20132806. https:\/\/doi.org\/10.1007\/s10489-018-01396-y","journal-title":"Appl Intell"},{"issue":"9","key":"1906_CR6","doi-asserted-by":"publisher","first-page":"5975","DOI":"10.1109\/TII.2020.2971057","volume":"16","author":"A Villalonga","year":"2020","unstructured":"Villalonga A, Beruvides G, Casta\u00f1o F, Haber RE (2020) Cloud-based industrial cyber\u2013physical system for data-driven reasoning: a review and use case on an industry 4.0 pilot line. IEEE Trans Ind Informatics 16(9):5975\u20135984. https:\/\/doi.org\/10.1109\/TII.2020.2971057","journal-title":"IEEE Trans Ind Informatics"},{"issue":"1","key":"1906_CR7","doi-asserted-by":"publisher","first-page":"13","DOI":"10.1109\/37.257890","volume":"14","author":"V Gullapalli","year":"1994","unstructured":"Gullapalli V, Franklin JA, Benbrahim H (1994) Acquiring robot SKILLS via reinforcement learning. IEEE Control Syst Mag 14(1):13\u201324. https:\/\/doi.org\/10.1109\/37.257890","journal-title":"IEEE Control Syst Mag"},{"issue":"4","key":"1906_CR8","doi-asserted-by":"publisher","first-page":"941","DOI":"10.1109\/72.508937","volume":"7","author":"BH Yang","year":"1996","unstructured":"Yang BH, Asada H (1996) Progressive learning and its application to robot impedance learning. IEEE Trans Neural Netw 7(4):941\u2013952. https:\/\/doi.org\/10.1109\/72.508937","journal-title":"IEEE Trans Neural Netw"},{"issue":"1","key":"1906_CR9","doi-asserted-by":"publisher","first-page":"101","DOI":"10.1016\/s0166-3615(97)00015-8","volume":"33","author":"M Nuttin","year":"1997","unstructured":"Nuttin M, VanBrussel H (1997) Learning the peg-into-hole assembly operation with a connectionist reinforcement technique. Comput Ind 33(1):101\u2013109. https:\/\/doi.org\/10.1016\/s0166-3615(97)00015-8","journal-title":"Comput Ind"},{"issue":"5","key":"1906_CR10","doi-asserted-by":"publisher","first-page":"1054","DOI":"10.1109\/TNN.1998.712192","volume":"9","author":"RS Sutton","year":"1998","unstructured":"Sutton RS, Barto AG (1998) Reinforcement learning: an introduction. IEEE Trans Neural Netw 9(5):1054\u20131054. https:\/\/doi.org\/10.1109\/TNN.1998.712192","journal-title":"IEEE Trans Neural Netw"},{"issue":"12","key":"1906_CR11","doi-asserted-by":"publisher","first-page":"4211","DOI":"10.1007\/s10489-019-01487-4","volume":"49","author":"S Ding","year":"2019","unstructured":"Ding S, Du W, Zhao X, Wang L, Jia W (2019) A new asynchronous reinforcement learning algorithm based on improved parallel PSO. Appl Intell 49(12):4211\u20134222. https:\/\/doi.org\/10.1007\/s10489-019-01487-4","journal-title":"Appl Intell"},{"issue":"10","key":"1906_CR12","doi-asserted-by":"publisher","first-page":"3749","DOI":"10.1007\/s10489-019-01484-7","volume":"49","author":"P Liu","year":"2019","unstructured":"Liu P, Zhao Y, Zhao W, Tang X, Yang Z (2019) An exploratory rollout policy for imagination-augmented agents. Appl Intell 49(10):3749\u20133764. https:\/\/doi.org\/10.1007\/s10489-019-01484-7","journal-title":"Appl Intell"},{"issue":"2","key":"1906_CR13","doi-asserted-by":"publisher","first-page":"1153","DOI":"10.1109\/tii.2018.2826064","volume":"15","author":"CG Yang","year":"2019","unstructured":"Yang CG, Zeng C, Cong Y, Wang N, Wang M (2019) A learning framework of adaptive manipulative Skills from human to robot. Ieee Trans Ind Informatics 15(2):1153\u20131161. https:\/\/doi.org\/10.1109\/tii.2018.2826064","journal-title":"Ieee Trans Ind Informatics"},{"issue":"4","key":"1906_CR14","doi-asserted-by":"publisher","first-page":"1600","DOI":"10.1109\/tmech.2017.2671342","volume":"22","author":"A Wan","year":"2017","unstructured":"Wan A, Xu J, Chen HP, Zhang S, Chen K (2017) Optimal path planning and control of assembly robots for hard-measuring easy-deformation assemblies. Ieee-Asme Trans Mechatron 22(4):1600\u20131609. https:\/\/doi.org\/10.1109\/tmech.2017.2671342","journal-title":"Ieee-Asme Trans Mechatron"},{"issue":"3","key":"1906_CR15","doi-asserted-by":"publisher","first-page":"1658","DOI":"10.1109\/tii.2018.2868859","volume":"15","author":"J Xu","year":"2019","unstructured":"Xu J, Hou ZM, Wang W, Xu BH, Zhang KG, Chen K (2019) Feedback deep deterministic policy gradient with fuzzy reward for robotic multiple peg-in-hole assembly tasks. Ieee Trans Ind Informatics 15(3):1658\u20131667. https:\/\/doi.org\/10.1109\/tii.2018.2868859","journal-title":"Ieee Trans Ind Informatics"},{"issue":"1","key":"1906_CR16","doi-asserted-by":"publisher","first-page":"118","DOI":"10.1109\/87.817697","volume":"8","author":"K Young Ho","year":"2000","unstructured":"Young Ho K, Lewis FL (2000) Reinforcement adaptive learning neural-net-based friction compensation control for high speed and precision. IEEE Trans Control Syst Technol 8(1):118\u2013126. https:\/\/doi.org\/10.1109\/87.817697","journal-title":"IEEE Trans Control Syst Technol"},{"issue":"2","key":"1906_CR17","doi-asserted-by":"publisher","first-page":"408","DOI":"10.1109\/TPAMI.2013.218","volume":"37","author":"MP Deisenroth","year":"2015","unstructured":"Deisenroth MP, Fox D, Rasmussen CE (2015) Gaussian processes for data-efficient learning in robotics and control. IEEE Trans Pattern Anal Mach Intell 37(2):408\u2013423. https:\/\/doi.org\/10.1109\/TPAMI.2013.218","journal-title":"IEEE Trans Pattern Anal Mach Intell"},{"key":"1906_CR18","unstructured":"Schulman J, Wolski F, Dhariwal P, Radford A, Klimov O (2017) Proximal policy optimization algorithms. arXiv e-prints"},{"key":"1906_CR19","volume-title":"Asynchronous methods for deep reinforcement learning","author":"V Mnih","year":"2016","unstructured":"Mnih V, Badia AP, Mirza M, Graves A, Harley T, Lillicrap TP, Silver D, Kavukcuoglu K (2016) Asynchronous methods for deep reinforcement learning, vol 48. Paper presented at the proceedings of the 33rd International Conference on International Conference on Machine Learning, New York"},{"key":"1906_CR20","unstructured":"Wang Z, Bapst V, Heess N, Mnih V, Munos R, Kavukcuoglu K, de Freitas N (2016) Sample efficient actor-critic with experience replay"},{"key":"1906_CR21","doi-asserted-by":"publisher","unstructured":"Mousavi SS, Schukat M, Howley E (2018) Deep reinforcement learning: an overview. In: Bi Y, Kapoor S, Bhatia R (eds) Proceedings of Sai Intelligent Systems Conference, vol 16. Lecture Notes in Networks and Systems. pp 426-440. https:\/\/doi.org\/10.1007\/978-3-319-56991-8_32","DOI":"10.1007\/978-3-319-56991-8_32"},{"issue":"2","key":"1906_CR22","doi-asserted-by":"publisher","first-page":"70","DOI":"10.1109\/tamd.2010.2051031","volume":"2","author":"S Singh","year":"2010","unstructured":"Singh S, Lewis RL, Barto AG, Sorg J (2010) Intrinsically motivated reinforcement learning: an evolutionary perspective. IEEE Trans Auton Ment Dev 2(2):70\u201382. https:\/\/doi.org\/10.1109\/tamd.2010.2051031","journal-title":"IEEE Trans Auton Ment Dev"},{"issue":"7540","key":"1906_CR23","doi-asserted-by":"publisher","first-page":"529","DOI":"10.1038\/nature14236","volume":"518","author":"V Mnih","year":"2015","unstructured":"Mnih V, Kavukcuoglu K, Silver D, Rusu AA, Veness J, Bellemare MG, Graves A, Riedmiller M, Fidjeland AK, Ostrovski G, Petersen S, Beattie C, Sadik A, Antonoglou I, King H, Kumaran D, Wierstra D, Legg S, Hassabis D (2015) Human-level control through deep reinforcement learning. Nature 518(7540):529\u2013533. https:\/\/doi.org\/10.1038\/nature14236","journal-title":"Nature"},{"key":"1906_CR24","doi-asserted-by":"publisher","first-page":"101","DOI":"10.1146\/annurev-psych-122414-033625","volume":"68","author":"SJ Gershman","year":"2017","unstructured":"Gershman SJ, Daw ND (2017) Reinforcement learning and episodic memory in humans and animals: an integrative framework. Annu Rev Psychol 68:101\u2013128. https:\/\/doi.org\/10.1146\/annurev-psych-122414-033625","journal-title":"Annu Rev Psychol"},{"issue":"1","key":"1906_CR25","doi-asserted-by":"publisher","first-page":"14","DOI":"10.1109\/TSMCB.2010.2043839","volume":"41","author":"FL Lewis","year":"2011","unstructured":"Lewis FL, Vamvoudakis KG (2011) Reinforcement learning for partially observable dynamic processes: adaptive dynamic programming using measured output data. IEEE Trans Syst Man Cybernetics Part B (Cybernetics) 41(1):14\u201325. https:\/\/doi.org\/10.1109\/TSMCB.2010.2043839","journal-title":"IEEE Trans Syst Man Cybernetics Part B (Cybernetics)"},{"key":"1906_CR26","unstructured":"Lillicrap TP, Hunt JJ, Pritzel A, Heess N, Erez T, Tassa Y, Silver D, Wierstra D (2015) Continuous control with deep reinforcement learning. Computer Science"},{"key":"1906_CR27","unstructured":"Haarnoja T, Zhou A, Abbeel P, Levine S (2018) Soft actor-critic: off-policy maximum entropy deep reinforcement learning with a stochastic actor. Paper presented at the Proceedings of the 35th International Conference on Machine Learning, Proceedings of Machine Learning Research,"},{"key":"1906_CR28","unstructured":"Schulman J, Levine S, Moritz P, Jordan MI, Abbeel P (2015) Trust region policy optimization. arXiv e-prints:arXiv:1502.05477"}],"container-title":["Applied Intelligence"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10489-020-01906-x.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s10489-020-01906-x\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s10489-020-01906-x.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2021,5,20]],"date-time":"2021-05-20T08:03:18Z","timestamp":1621497798000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s10489-020-01906-x"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2020,11,13]]},"references-count":28,"journal-issue":{"issue":"6","published-print":{"date-parts":[[2021,6]]}},"alternative-id":["1906"],"URL":"https:\/\/doi.org\/10.1007\/s10489-020-01906-x","relation":{},"ISSN":["0924-669X","1573-7497"],"issn-type":[{"value":"0924-669X","type":"print"},{"value":"1573-7497","type":"electronic"}],"subject":[],"published":{"date-parts":[[2020,11,13]]},"assertion":[{"value":"13 November 2020","order":1,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Compliance with ethical standards"}},{"value":"The authors declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declaration of interests"}}]}}