{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,4,9]],"date-time":"2026-04-09T11:00:19Z","timestamp":1775732419714,"version":"3.50.1"},"reference-count":47,"publisher":"Springer Science and Business Media LLC","issue":"1","license":[{"start":{"date-parts":[[2021,1,5]],"date-time":"2021-01-05T00:00:00Z","timestamp":1609804800000},"content-version":"tdm","delay-in-days":0,"URL":"http:\/\/www.springer.com\/tdm"},{"start":{"date-parts":[[2021,1,5]],"date-time":"2021-01-05T00:00:00Z","timestamp":1609804800000},"content-version":"vor","delay-in-days":0,"URL":"http:\/\/www.springer.com\/tdm"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Prog Artif Intell"],"published-print":{"date-parts":[[2021,3]]},"DOI":"10.1007\/s13748-020-00225-z","type":"journal-article","created":{"date-parts":[[2021,1,5]],"date-time":"2021-01-05T02:03:01Z","timestamp":1609812181000},"page":"83-97","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":33,"title":["A synchronous deep reinforcement learning model for automated multi-stock trading"],"prefix":"10.1007","volume":"10","author":[{"ORCID":"https:\/\/orcid.org\/0000-0002-1654-1229","authenticated-orcid":false,"given":"Rasha","family":"AbdelKawy","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-8726-605X","authenticated-orcid":false,"given":"Walid M.","family":"Abdelmoez","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0002-8024-3795","authenticated-orcid":false,"given":"Amin","family":"Shoukry","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2021,1,5]]},"reference":[{"key":"225_CR1","doi-asserted-by":"publisher","unstructured":"Hasbrouck., J.: 22 Modeling market microstructure time series, In: Handbook of Statistics, Vol. 14, pp. 647-692, ELSEVIER(1996). https:\/\/doi.org\/10.1016\/S0169-7161(96)14024-4","DOI":"10.1016\/S0169-7161(96)14024-4"},{"issue":"1","key":"225_CR2","doi-asserted-by":"publisher","first-page":"259","DOI":"10.1016\/j.eswa.2014.07.040","volume":"42","author":"J Pate","year":"2015","unstructured":"Pate, J., Shah, S., Thakkar, P.: Predicting stock and stock price index movement using trend deterministic data preparation and machine learning techniques. Exp. Syst. Appl. 42(1), 259\u2013268 (2015). https:\/\/doi.org\/10.1016\/j.eswa.2014.07.040. Elsevier","journal-title":"Exp. Syst. Appl."},{"key":"225_CR3","doi-asserted-by":"publisher","first-page":"194","DOI":"10.1016\/j.eswa.2016.02.006","volume":"55","author":"RC Cavalcantea","year":"2016","unstructured":"Cavalcantea, R.C., Brasileirob, R.C., Souza, V.L., Nobrega, J.P., Oliveirab, A.L.I.: Computational intelligence and financial markets: A survey and future directions. Exp. Syst. Appl. 55, 194\u2013211 (2016). https:\/\/doi.org\/10.1016\/j.eswa.2016.02.006","journal-title":"Exp. Syst. Appl."},{"issue":"2","key":"225_CR4","doi-asserted-by":"publisher","first-page":"297","DOI":"10.1007\/s10898-011-9692-3","volume":"53","author":"P Gupta","year":"2012","unstructured":"Gupta, P., Mehlawat, M.K., Mittal, G.: Asset portfolio optimization using support vector machines and real-coded genetic algorithm. J. Global Optim. 53(2), 297\u2013315 (2012)","journal-title":"J. Global Optim."},{"key":"225_CR5","doi-asserted-by":"publisher","unstructured":"Yang, B., Gong, Z.-J., Yang, W.: Stock market index prediction using deep neural network ensemble, In: 36th Chinese Control Conference (CCC), pp. 26-28, Dalian, China (2017). https:\/\/doi.org\/10.23919\/ChiCC.2017.8027964","DOI":"10.23919\/ChiCC.2017.8027964"},{"key":"225_CR6","doi-asserted-by":"publisher","first-page":"60","DOI":"10.1016\/j.eswa.2017.12.026","volume":"97","author":"J Zhang","year":"2018","unstructured":"Zhang, J., Shicheng, C., Yan, X., Qianmu, L., Tao, L.: A novel data-driven stock price trend prediction system. Exp. Syst. Appl. 97, 60\u201369 (2018). https:\/\/doi.org\/10.1016\/j.eswa.2017.12.026","journal-title":"Exp. Syst. Appl."},{"key":"225_CR7","doi-asserted-by":"publisher","first-page":"187","DOI":"10.1016\/j.eswa.2017.04.030","volume":"83","author":"E Chonga","year":"2017","unstructured":"Chonga, E., Han, C., Parka, F.C.: Deep learning networks for stock market analysis and prediction: Methodology, data representations, and case studies. Exp. Syst. Appl. 83, 187\u2013205 (2017)","journal-title":"Exp. Syst. Appl."},{"issue":"4","key":"225_CR8","doi-asserted-by":"publisher","first-page":"e0230635","DOI":"10.1371\/journal.pone.0230635","volume":"V15","author":"J Lee","year":"2020","unstructured":"Lee, J., Kang, J.: Effectively training neural networks for stock index prediction: Predicting the S&P 500 index without using its index data. PLoS ONE V15(4), e0230635 (2020). https:\/\/doi.org\/10.1371\/journal.pone.0230635","journal-title":"PLoS ONE"},{"key":"225_CR9","doi-asserted-by":"publisher","unstructured":"Sezer, O., Ozbayoglu, M.: Algorithmic financial trading with deep convolutional neural networks: Time series to image conversion approach. Appl. Soft Comput. 70 (2018) https:\/\/doi.org\/10.1016\/j.asoc.2018.04.024","DOI":"10.1016\/j.asoc.2018.04.024"},{"key":"225_CR10","doi-asserted-by":"crossref","unstructured":"Jiang, W.: Applications of deep learning in stock market prediction: recent progress, arXiv:2003.01859 (2020), Preprint submitted to Elsevier Journal","DOI":"10.1016\/j.eswa.2021.115537"},{"key":"225_CR11","doi-asserted-by":"publisher","first-page":"106384","DOI":"10.1016\/j.asoc.2020.106384","volume":"93","author":"AM Murat","year":"2020","unstructured":"Murat, A.M., Omer, M.U., Sezer, B.S.: Deep learning for financial applications : A survey. Appl. Soft Comput. 93, 106384 (2020). https:\/\/doi.org\/10.1016\/j.asoc.2020.106384","journal-title":"Appl. Soft Comput."},{"key":"225_CR12","first-page":"354","volume":"550","author":"D Silver","year":"2017","unstructured":"Silver, D., Schrittwieser, J., Simonyan, K., Antonoglou, I., Huang, A., Guez, A., Hubert, T., Baker, L., Lai, M., Bolton, A., Chen, Y., Lillicrap, T., Hui, F., Sifre, L., Driessche, G., Graepel, T., Hassabis, D.: Mastering the game of go without human knowledge. Int. J. Sci. Nat. 550, 354\u2013359 (2017)","journal-title":"Int. J. Sci. Nat."},{"key":"225_CR13","unstructured":"https:\/\/colah.github.io\/posts\/2015-08-Understanding-LSTMs\/ Accessed (August 2020)"},{"key":"225_CR14","first-page":"279","volume-title":"Q-learning, Machine Learning","author":"CJ Watkins","year":"1992","unstructured":"Watkins, C.J., Dayan, P.: Q-learning, Machine Learning, vol. 8, pp. 279\u2013292. Springer, Berlin (1992)"},{"key":"225_CR15","doi-asserted-by":"publisher","first-page":"529","DOI":"10.1038\/nature14236","volume":"518","author":"V Mnih","year":"2015","unstructured":"Mnih, V., Kavukcuoglu, K., Silver, D., Rusu, A.A., Veness, J., Bellemare, M.G., Graves, A., Riedmiller, M., Fidjeland, A.K., Ostrovski, G., Petersen, S., Beattie, C., Sadik, A., Antonoglou, I., King, H., Kumaran, D., Wierstra, D., Legg, S., Hassabis, D.: Human-level control through deep reinforcement learning. Nature 518, 529\u2013533 (2015)","journal-title":"Nature"},{"key":"225_CR16","doi-asserted-by":"publisher","first-page":"85","DOI":"10.1016\/j.neunet.2014.09.003","volume":"61","author":"Huber J Schmid","year":"2015","unstructured":"Schmid, Huber J.: Deep learning in neural networks: An overview. Neural Netw. V 61, 85\u2013117 (2015)","journal-title":"Neural Netw. V"},{"issue":"7","key":"225_CR17","doi-asserted-by":"publisher","first-page":"1527","DOI":"10.1162\/neco.2006.18.7.1527","volume":"V18","author":"GE Hinton","year":"2006","unstructured":"Hinton, G.E., Osindero, S., Teh, Y.W.: A fast learning algorithm for deep belief nets. Neural Comput. V18(7), 1527\u20131554 (2006)","journal-title":"Neural Comput."},{"issue":"19","key":"225_CR18","first-page":"153","volume":"V","author":"Y Bengio","year":"2006","unstructured":"Bengio, Y., Lamblin, P., Popovici, D., Larochelle, H.: Greedy layer-wise training of deep networks. Adv. Neural Inf. Process. Syst. V(19), 153\u2013160 (2006)","journal-title":"Adv. Neural Inf. Process. Syst."},{"key":"225_CR19","first-page":"2613","volume":"23","author":"HV Hasselt","year":"2010","unstructured":"Hasselt, H.V.: Double Q-learning. Adv. Neural Inf. Process. Syst. 23, 2613\u20132621 (2010)","journal-title":"Adv. Neural Inf. Process. Syst."},{"key":"225_CR20","unstructured":"Wang, Z., Freitas, N., de., Lanctot, M.: Dueling network architectures for deep reinforcement learning, In the International Conference on Machine Learning (ICML), (2015). arXiv preprint arXiv:1511.06581"},{"key":"225_CR21","unstructured":"Hessel, M., Modayil, J., van Hasselt, H., Schaul, T., Ostrovski, G., Dabney, W., Horgan, D., Piot, B., Azar, M., Silver, D.: Rainbow: Combining Improvements in Deep Reinforcement Learning, Thirty-Second AAAI Conference on Artificial Intelligence (2017). arXiv preprint arXiv:1710.02298"},{"key":"225_CR22","unstructured":"Sutton, R.S., McAllester, D.A., Singh, S.P., Mansour, Y.: Policy gradient methods for reinforcement learning with function approximation, Advances in Neural Information Processing Systems, Vol. 12, pp. 1057\u20131063. (NIPS 1999) MIT Press, Cambridge, MA (2000)"},{"key":"225_CR23","unstructured":"Lillicrap, T.P., Hunt, J.J., Pritzel, A., Heess, N., Erez, T., Tassa, Y., Silver, D., Wierstra, D.: Continuous control with deep reinforcement learning, In: International Conference Learning Representations (2016). arXiv preprint arXiv:1509.02971"},{"key":"225_CR24","unstructured":"Schulman, J., Levine, S., Abbeel, P., Jordan, M.I., Moritz, P.: Trust Region Policy Optimization, In: 32nd International Conference on Machine Learning, Vol. 37, pp. 1889\u20131897, PMLR. http:\/\/proceedings.mlr.press\/v37\/schulman15.html(2015)"},{"key":"225_CR25","unstructured":"Schulman, J., Wolski, F., Dhariwal, P., Radford, A., Klimov, O.: Proximal policy optimization algorithms, Computing Research Repository (CoRR), 1707.06347 (2017). arXiv preprint arXiv:1707.06347"},{"key":"225_CR26","unstructured":"Wang, Z., Bapst, V., Heess, N., Mnih, V., Munos, R., Kavukcuoglu, K., Freitas, N.: Sample efficient actor-critic with experience replay,ICLR (2016). arXiv preprint arXiv:1611.01224"},{"key":"225_CR27","unstructured":"OpenAI, https:\/\/openai.com\/, Accessed 1.7 April 2020"},{"key":"225_CR28","unstructured":"OpenAI Baselines: ACKTR & A2C, https:\/\/openai.com\/blog\/baselines-acktr-a2c\/, Accessed 17 April 2020"},{"key":"225_CR29","unstructured":"Mnih, V., Kavukcuoglu, K., Silver, D., Graves, A., Antonoglou, I., Wierstra, D., Riedmil-ler, M.: Playing atari with deep reinforcement learning, In NIPS Deep Learning Work-shop (2013)"},{"key":"225_CR30","unstructured":"Mnih, V., Badia, A.P., Mirza, M., Graves, A., Lillicrap, T.P., Harley, T., Silver, D., Kavukcuoglu, K.: Asynchronous methods for deep reinforcement learning, In: 33rd International Conference on Machine Learning, Vol. 48, pp. 1928-1937, PMLR (2016)"},{"issue":"4","key":"225_CR31","doi-asserted-by":"publisher","first-page":"875","DOI":"10.1109\/72.935097","volume":"12","author":"J Moody","year":"2001","unstructured":"Moody, J., Saffell, M.: Learning to trade via direct reinforcement. IEEE Trans. Neural Netw. 12(4), 875\u2013889 (2001). https:\/\/doi.org\/10.1109\/72.935097","journal-title":"IEEE Trans. Neural Netw."},{"issue":"3","key":"225_CR32","doi-asserted-by":"publisher","first-page":"653","DOI":"10.1109\/TNNLS.2016.2522401","volume":"28","author":"Y Deng","year":"2017","unstructured":"Deng, Y., Bao, F., Youyong, K., Zhiquan, R., Qionghai, D.: Deep direct reinforcement learning for financial signal representation and trading. IEEE Trans. Neural Netw. Learn. Syst. 28(3), 653\u2013664 (2017)","journal-title":"IEEE Trans. Neural Netw. Learn. Syst."},{"issue":"87","key":"225_CR33","doi-asserted-by":"publisher","first-page":"267","DOI":"10.1016\/j.eswa.2017.06.023","volume":"V","author":"S Almahdi","year":"2017","unstructured":"Almahdi, S., Yang, S.Y.: An adaptive portfolio trading system: A risk-return portfolio optimization using recurrent reinforcement learning with expected maximum drawdown. Exp. Syst. Appl. V(87), 267\u2013279 (2017). https:\/\/doi.org\/10.1016\/j.eswa.2017.06.023","journal-title":"Exp. Syst. Appl."},{"key":"225_CR34","unstructured":"Jiang, Z., Xu, D., Liang, J.: A deep reinforcement learning framework for the financial portfolio management problem, arXiv:1706.10059 (2017)"},{"key":"225_CR35","unstructured":"https:\/\/github.com\/OLPS\/OLPS, last accessed October 2020"},{"key":"225_CR36","unstructured":"Li, B., Sahoo, D., S. CH. Hoi.: Olps: A toolbox for online portfolio selection., J. Mach. Learn. Res. (JMLR), (2015)"},{"issue":"3","key":"225_CR37","first-page":"35","volume":"V46","author":"B Li","year":"2014","unstructured":"Li, B., Hoi, S.C.H.: Online portfolio selection: A survey. ACM Comput. Surv. (CSUR) V46(3), 35 (2014)","journal-title":"ACM Comput. Surv. (CSUR)"},{"key":"225_CR38","unstructured":"https:\/\/github.com\/ZhengyaoJiang\/PGPortfolio , last accessed October 2020"},{"key":"225_CR39","doi-asserted-by":"crossref","unstructured":"Jiang, Z., Liang, J.: Cryptocurrency portfolio management with deep reinforcement learning., Intelligent Systems Conference., SAI Conferences,2017. Preprint: arXiv:1612.01277","DOI":"10.1109\/IntelliSys.2017.8324237"},{"key":"225_CR40","unstructured":"Liang, Z., Chen, H., Zhu, J., Jiang, K., Li, Y.: Adversarial Deep Reinforcement Learning in Portfolio Management, arXiv:1808.09940 (2018)"},{"key":"225_CR41","doi-asserted-by":"crossref","unstructured":"Hegde, S., Kumar, V., Singh, A.: Risk aware portfolio construction using deep deterministic policy gradients, IEEE Symposium Series on Computational Intelligence (2018)","DOI":"10.1109\/SSCI.2018.8628791"},{"key":"225_CR42","doi-asserted-by":"crossref","unstructured":"Wang, J., Zhang, Y., Tang, K., Wu, J., Xiong, Z.: AlphaStock: A Buying Winners-and-Selling-Losers Investment Strategy using Interpretable Deep Reinforcement Attention Networks, 25th ACM SIGKDD, pp.1900-1908 (2019)","DOI":"10.1145\/3292500.3330647"},{"key":"225_CR43","doi-asserted-by":"crossref","unstructured":"Li, Y., Zheng, W., Zheng, Z.: Deep Robust Reinforcement Learning for Practical Algorithmic Trading, IEEE Access, pp.108014\u2013108022 (2019)","DOI":"10.1109\/ACCESS.2019.2932789"},{"key":"225_CR44","doi-asserted-by":"publisher","first-page":"113456","DOI":"10.1016\/j.eswa.2020.113456","volume":"156","author":"F Soleymani","year":"2020","unstructured":"Soleymani, F., Elodie, P.: Financial portfolio optimization with online deep reinforcement learning and restricted stacked autoencoder - DeepBreath\u201d. Exp. Syst. Appl. 156, 113456 (2020)","journal-title":"Exp. Syst. Appl."},{"issue":"7","key":"225_CR45","doi-asserted-by":"publisher","first-page":"e0236178","DOI":"10.1371\/journal.pone.0236178","volume":"15","author":"J Leem","year":"2020","unstructured":"Leem, J., Kim, H.Y.: Action specialized expert ensemble trading system with extended discrete action space using deep reinforcement learning. PLoS ONE 15(7), e0236178 (2020). https:\/\/doi.org\/10.1371\/journal.pone.0236178","journal-title":"PLoS ONE"},{"key":"225_CR46","doi-asserted-by":"publisher","unstructured":"Mosavi, A., Ghamisi, P., Faghan, Y., Duan, P.: Shamshirband. Comprehensive Review of Deep Reinforcement Learning Methods and Applications in Economics. (2020). https:\/\/doi.org\/10.20944\/preprints202003.0309.v1","DOI":"10.20944\/preprints202003.0309.v1"},{"key":"225_CR47","unstructured":"Charpentier, A., Elie, R., Remlinger, C.: Reinforcement Learning in Economics and Finance (2020) arXiv:2003.10014"}],"container-title":["Progress in Artificial Intelligence"],"original-title":[],"language":"en","link":[{"URL":"http:\/\/link.springer.com\/content\/pdf\/10.1007\/s13748-020-00225-z.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"http:\/\/link.springer.com\/article\/10.1007\/s13748-020-00225-z\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"http:\/\/link.springer.com\/content\/pdf\/10.1007\/s13748-020-00225-z.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2022,12,10]],"date-time":"2022-12-10T14:08:25Z","timestamp":1670681305000},"score":1,"resource":{"primary":{"URL":"http:\/\/link.springer.com\/10.1007\/s13748-020-00225-z"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2021,1,5]]},"references-count":47,"journal-issue":{"issue":"1","published-print":{"date-parts":[[2021,3]]}},"alternative-id":["225"],"URL":"https:\/\/doi.org\/10.1007\/s13748-020-00225-z","relation":{},"ISSN":["2192-6352","2192-6360"],"issn-type":[{"value":"2192-6352","type":"print"},{"value":"2192-6360","type":"electronic"}],"subject":[],"published":{"date-parts":[[2021,1,5]]},"assertion":[{"value":"28 April 2020","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"16 November 2020","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"5 January 2021","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}}]}}