{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,5,13]],"date-time":"2026-05-13T16:31:05Z","timestamp":1778689865790,"version":"3.51.4"},"reference-count":48,"publisher":"Springer Science and Business Media LLC","issue":"3","license":[{"start":{"date-parts":[[2023,4,8]],"date-time":"2023-04-08T00:00:00Z","timestamp":1680912000000},"content-version":"tdm","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"},{"start":{"date-parts":[[2023,4,8]],"date-time":"2023-04-08T00:00:00Z","timestamp":1680912000000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.springernature.com\/gp\/researchers\/text-and-data-mining"}],"content-domain":{"domain":["link.springer.com"],"crossmark-restriction":false},"short-container-title":["Evolving Systems"],"published-print":{"date-parts":[[2024,6]]},"DOI":"10.1007\/s12530-023-09498-w","type":"journal-article","created":{"date-parts":[[2023,4,8]],"date-time":"2023-04-08T09:02:40Z","timestamp":1680944560000},"page":"717-729","update-policy":"https:\/\/doi.org\/10.1007\/springer_crossmark_policy","source":"Crossref","is-referenced-by-count":8,"title":["Creating an AI fashioner through deep learning and computer vision"],"prefix":"10.1007","volume":"15","author":[{"ORCID":"https:\/\/orcid.org\/0000-0002-1010-129X","authenticated-orcid":false,"given":"Caner","family":"Balim","sequence":"first","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]},{"ORCID":"https:\/\/orcid.org\/0000-0003-2252-2128","authenticated-orcid":false,"given":"Kemal","family":"Ozkan","sequence":"additional","affiliation":[],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"297","published-online":{"date-parts":[[2023,4,8]]},"reference":[{"key":"9498_CR1","doi-asserted-by":"crossref","unstructured":"Anderson P, He X, Buehler C, Teney D, Johnson M, Gould S, Zhang L (2017) Bottom-up and top-down attention for image captioning and visual question answering. In: Proceedings of the IEEE computer society conference on computer vision and pattern recognition, pp 6077\u20136086","DOI":"10.1109\/CVPR.2018.00636"},{"key":"9498_CR2","doi-asserted-by":"publisher","first-page":"614","DOI":"10.46519\/ij3dptdi.991789","volume":"5","author":"C Balim","year":"2021","unstructured":"Balim C, \u00d6zkan K (2021) Ur\u00fcn g\u00f6rsellerini kullanarak e-ticaret sistemleri i\u00e7in \u00fcr\u00fcn ba\u015fli\u011fi olu\u015fturulmasi. Int J 3D Rint Technol Dig Ind 5:614\u2013624. https:\/\/doi.org\/10.46519\/ij3dptdi.991789","journal-title":"Int J 3D Rint Technol Dig Ind"},{"key":"9498_CR3","doi-asserted-by":"publisher","first-page":"119305","DOI":"10.1016\/j.eswa.2022.119305","volume":"215","author":"C Balim","year":"2023","unstructured":"Balim C, \u00d6zkan K (2023) Diagnosing fashion outfit compatibility with deep learning techniques. Expert Syst Appl 215:119305. https:\/\/doi.org\/10.1016\/j.eswa.2022.119305","journal-title":"Expert Syst Appl"},{"key":"9498_CR4","unstructured":"Banerjee S, Lavie A (2005) METEOR: an automatic metric for MT evaluation with improved correlation with human judgments. In: Proceedings of the ACL workshop on intrinsic and extrinsic evaluation measures for machine translation and\/or summarization, pp 65\u201372"},{"key":"9498_CR5","doi-asserted-by":"crossref","unstructured":"Chen L, He Y (2018) Dress fashionably: learn fashion collocation with deep mixed-category metric learning. In: Proceedings of the AAAI Conference on Artificial Intelligence, vol 32, no. 1","DOI":"10.1609\/aaai.v32i1.11895"},{"key":"9498_CR6","doi-asserted-by":"publisher","unstructured":"Chen X, Chen H, Xu H, Zhang Y, Cao Y, Qin Z, Zha H (2019a) Personalized fashion recommendation with visual explanations based on multimodal attention network: towards visually explainable recommendation. In: Proceedings of the 42nd International ACM SIGIR conference on research and development in information retrieval, pp 765\u2013774. Association for Computing Machinery, New York. https:\/\/doi.org\/10.1145\/3331184.3331254","DOI":"10.1145\/3331184.3331254"},{"key":"9498_CR7","doi-asserted-by":"crossref","unstructured":"Chen W, Huang P, Xu J, Guo X, Guo C, Sun F, Li C, Pfadler A, Zhao H, Zhao B (2019b) POG: personalized outfit generation for fashion recommendation at alibaba iFashion. In: Proceedings of the ACM SIGKDD international conference on knowledge discovery and data mining, pp 2662\u20132670","DOI":"10.1145\/3292500.3330652"},{"key":"9498_CR8","doi-asserted-by":"crossref","unstructured":"Cho K, van Merrienboer B, Gulcehre C, Bahdanau D, Bougares F, Schwenk H, Bengio Y (2014) Learning phrase representations using RNN encoder-decoder for statistical machine translation. arXiv:1406.1078 [cs, stat]","DOI":"10.3115\/v1\/D14-1179"},{"key":"9498_CR9","unstructured":"FashionVLP (2022) Vision language transformer for fashion retrieval with feedback, https:\/\/www.amazon.science\/publications\/fashionvlp-vision-language-transformer-for-fashion-retrieval-with-feedback. Accessed 8 Aug 2022"},{"key":"9498_CR10","doi-asserted-by":"publisher","unstructured":"Han X, Wu Z, Jiang Y-G, Davis LS (2017) Learning fashion compatibility with bidirectional LSTMs. In: MM 2017\u2014proceedings of the 2017 ACM multimedia conference, pp 1078\u20131086. Doi: https:\/\/doi.org\/10.1145\/3123266.3123394","DOI":"10.1145\/3123266.3123394"},{"key":"9498_CR11","doi-asserted-by":"publisher","unstructured":"Han X (2022) Prototype-guided Attribute-wise Interpretable Scheme for Clothing Matching. In: Proceedings of the 42nd International ACM SIGIR conference on research and development in information retrieval. https:\/\/doi.org\/10.1145\/3331184.3331245. Accessed 7 Aug 2022","DOI":"10.1145\/3331184.3331245"},{"key":"9498_CR12","doi-asserted-by":"crossref","unstructured":"He R, Packer C, McAuley J (2016) Learning compatibility across categories for heterogeneous item recommendation. In: Proceedings\u2014IEEE international conference on data mining, ICDM, pp 937\u2013942","DOI":"10.1109\/ICDM.2016.0116"},{"key":"9498_CR13","doi-asserted-by":"crossref","unstructured":"He K, Gkioxari G, Doll\u00e1r P, Girshick R (2017) Mask r-cnn. In: Proceedings of the IEEE international conference on computer vision, pp 2961\u20132969","DOI":"10.1109\/ICCV.2017.322"},{"key":"9498_CR14","unstructured":"Herdade S, Kappeler A, Boakye K, Soares J (2019) Image captioning: transforming objects into words. Advances in neural information processing systems 32"},{"key":"9498_CR15","unstructured":"Ji Y-H, Jun H, Kim I, Kim J, Kim Y, Ko B, Kook H-K, Lee J, Lee S, Park S (2020) An effective pipeline for a real-world clothes retrieval system. arXiv:2005.12739 [cs]"},{"key":"9498_CR16","doi-asserted-by":"publisher","unstructured":"Kaicheng P, Xingxing Z, Wong WK (2021) modeling fashion compatibility with explanation by using bidirectional LSTM. In: 2021 IEEE\/CVF conference on computer vision and pattern recognition workshops (CVPRW), pp 3889\u20133893. https:\/\/doi.org\/10.1109\/CVPRW53098.2021.00432","DOI":"10.1109\/CVPRW53098.2021.00432"},{"key":"9498_CR17","doi-asserted-by":"publisher","first-page":"1833","DOI":"10.1109\/TCYB.2018.2887094","volume":"50","author":"Z Kang","year":"2020","unstructured":"Kang Z, Pan H, Hoi SCH, Xu Z (2020a) Robust graph learning from noisy data. IEEE Trans Cybern 50:1833\u20131843. https:\/\/doi.org\/10.1109\/TCYB.2018.2887094","journal-title":"IEEE Trans Cybern"},{"key":"9498_CR18","doi-asserted-by":"crossref","unstructured":"Kang Z, Lu X, Liang J, Bai K, Xu Z (2020b) Relation-guided representation learning. arXiv:2007.05742 [cs, stat]","DOI":"10.1016\/j.neunet.2020.07.014"},{"key":"9498_CR19","doi-asserted-by":"publisher","DOI":"10.1016\/j.matpr.2020.09.365","author":"K Kavitha","year":"2020","unstructured":"Kavitha K, Kumar SL, Pravalika P, Sruthi K, Lalitha RVS, Rao NVK (2020) Fashion compatibility using convolutional neural networks. Mater Today: Proc. https:\/\/doi.org\/10.1016\/j.matpr.2020.09.365","journal-title":"Mater Today: Proc"},{"key":"9498_CR20","doi-asserted-by":"publisher","first-page":"1946","DOI":"10.1109\/TMM.2017.2690144","volume":"19","author":"Y Li","year":"2016","unstructured":"Li Y, Cao L, Zhu J, Luo J (2016) Mining fashion outfit composition using an end-to-end deep learning approach on set data. IEEE Trans Multimedia 19:1946\u20131955. https:\/\/doi.org\/10.1109\/TMM.2017.2690144","journal-title":"IEEE Trans Multimedia"},{"key":"9498_CR21","doi-asserted-by":"publisher","first-page":"68","DOI":"10.1016\/j.patrec.2020.12.001","volume":"141","author":"X Li","year":"2021","unstructured":"Li X, Ye Z, Zhang Z, Zhao M (2021) Clothes image caption generation with attribute detection and visual attention model. Pattern Recogn Lett 141:68\u201374. https:\/\/doi.org\/10.1016\/j.patrec.2020.12.001","journal-title":"Pattern Recogn Lett"},{"key":"9498_CR22","unstructured":"Li K, Liu C, Kumar R, Forsyth D (2019) Using discriminative methods to learn fashion compatibility across datasets. J Environ Sci (China) (English Ed)"},{"key":"9498_CR23","unstructured":"Lin CY (2004) Rouge: a package for automatic evaluation of summaries. In: Text summarization branches out, pp 74\u201381"},{"key":"9498_CR24","doi-asserted-by":"publisher","first-page":"1502","DOI":"10.1109\/TKDE.2019.2906190","volume":"32","author":"Y Lin","year":"2020","unstructured":"Lin Y, Ren P, Chen Z, Ren Z, Ma J, de Rijke M (2020) Explainable outfit recommendation with joint outfit matching and comment generation. IEEE Trans Knowl Data Eng 32:1502\u20131516. https:\/\/doi.org\/10.1109\/TKDE.2019.2906190","journal-title":"IEEE Trans Knowl Data Eng"},{"key":"9498_CR25","doi-asserted-by":"crossref","unstructured":"Liu Z, Luo P, Qiu S, Wang X, Tang X (2016) DeepFashion: powering robust clothes recognition and retrieval with rich annotations. In: Proceedings of IEEE conference on computer vision and pattern recognition (CVPR)","DOI":"10.1109\/CVPR.2016.124"},{"key":"9498_CR26","doi-asserted-by":"publisher","first-page":"150","DOI":"10.1016\/j.patrec.2021.04.009","volume":"147","author":"S Lu","year":"2021","unstructured":"Lu S, Zhu X, Wu Y, Wan X, Gao F (2021) Outfit compatibility prediction with multi-layered feature fusion network. Pattern Recogn Lett 147:150\u2013156. https:\/\/doi.org\/10.1016\/j.patrec.2021.04.009","journal-title":"Pattern Recogn Lett"},{"key":"9498_CR27","doi-asserted-by":"crossref","unstructured":"McAuley J, Targett C, Shi Q, Hengel A (2015) van den: image-based recommendations on styles and substitutes. In: SIGIR 2015\u2014 Proceedings of the 38th International ACM SIGIR conference on research and development in information retrieval, pp 43\u201352","DOI":"10.1145\/2766462.2767755"},{"key":"9498_CR28","doi-asserted-by":"publisher","first-page":"117333","DOI":"10.1016\/j.eswa.2022.117333","volume":"203","author":"D Mo","year":"2022","unstructured":"Mo D, Zou X, Wong W (2022) Neural stylist: towards online styling service. Expert Syst Appl 203:117333. https:\/\/doi.org\/10.1016\/j.eswa.2022.117333","journal-title":"Expert Syst Appl"},{"key":"9498_CR29","doi-asserted-by":"publisher","first-page":"311","DOI":"10.3115\/1073083.1073135","volume":"2011","author":"K Papineni","year":"2001","unstructured":"Papineni K, Roukos S, Ward T, Zhu W-J (2001) BLEU: a method for automatic evaluation of machine translation. ACL 2011:311\u2013318. https:\/\/doi.org\/10.3115\/1073083.1073135","journal-title":"ACL"},{"key":"9498_CR30","doi-asserted-by":"publisher","first-page":"138","DOI":"10.5392\/JKCA.2022.22.01.138","volume":"22","author":"YJ Park","year":"2022","unstructured":"Park YJ, Jo BC, Lee KU, Kim KS (2022) Improved transformer model for multimodal fashion recommendation conversation system. J Korea Contents Assoc 22:138\u2013147. https:\/\/doi.org\/10.5392\/JKCA.2022.22.01.138","journal-title":"J Korea Contents Assoc"},{"key":"9498_CR31","unstructured":"Qu W (2022) Visual and textual jointly enhanced interpretable fashion recommendation|IEEE Journals & Magazine|IEEE Xplore. https:\/\/ieeexplore.ieee.org\/document\/9046774. Accessed 7 Aug 2022."},{"key":"9498_CR32","doi-asserted-by":"crossref","unstructured":"Ren S, He K, Girshick R, Sun J (2016) Faster R-CNN: towards real-time object detection with region proposal networks. arXiv:1506.01497 [cs]","DOI":"10.1109\/TPAMI.2016.2577031"},{"key":"9498_CR33","doi-asserted-by":"crossref","unstructured":"Sidnev A, Krapivin A, Trushkov A, Krasikova E, Kazakov M, Viryasov M (2021) DeepMark++: real-time clothing detection at the edge. In: Presented at the proceedings of the IEEE\/CVF winter conference on applications of computer vision","DOI":"10.1109\/WACV48630.2021.00302"},{"key":"9498_CR34","doi-asserted-by":"publisher","unstructured":"Song X, Feng F, Liu J, Li Z, Nie L, Ma J (2017) NeuroStylist: neural compatibility modeling for clothing matching. In: Presented at the October 23. https:\/\/doi.org\/10.1145\/3123266.3123314","DOI":"10.1145\/3123266.3123314"},{"key":"9498_CR35","doi-asserted-by":"publisher","first-page":"237","DOI":"10.1016\/j.neucom.2018.06.098","volume":"395","author":"GL Sun","year":"2020","unstructured":"Sun GL, He JY, Wu X, Zhao B, Peng Q (2020a) Learning fashion compatibility across categories with deep multimodal neural networks. Neurocomputing 395:237\u2013246. https:\/\/doi.org\/10.1016\/j.neucom.2018.06.098","journal-title":"Neurocomputing"},{"key":"9498_CR36","doi-asserted-by":"crossref","unstructured":"Sun P, Wu L, Zhang K, Fu Y, Hong R, Wang M (2020b) Dual learning for explainable recommendation: towards unifying user preference prediction and review generation. In: Proceedings of the web conference 2020b, pp 837\u2013847. Association for Computing Machinery, New York","DOI":"10.1145\/3366423.3380164"},{"key":"9498_CR37","unstructured":"Sutskever I, Vinyals O, Le QV (2014) Sequence to sequence learning with neural networks. Advances in neural information processing systems 27"},{"key":"9498_CR38","doi-asserted-by":"publisher","DOI":"10.48550\/arXiv.1901.04870","author":"P Tangseng","year":"2019","unstructured":"Tangseng P, Okatani T (2019) Toward explainable fashion recommendation. Arxiv. https:\/\/doi.org\/10.48550\/arXiv.1901.04870","journal-title":"Arxiv"},{"key":"9498_CR39","doi-asserted-by":"crossref","unstructured":"Vasileva MI, Plummer BA, Dusad K, Rajpal S, Kumar R, Forsyth D (2018) Learning type-aware embeddings for fashion compatibility. In: Lecture notes in computer science (including subseries lecture notes in artificial intelligence and lecture notes in bioinformatics), 11220 LNCS, pp 405\u2013421","DOI":"10.1007\/978-3-030-01270-0_24"},{"key":"9498_CR40","unstructured":"Vaswani A, Shazeer N, Parmar N, Uszkoreit J, Jones L, Gomez AN, Kaiser \u0141, Polosukhin I (2017) Attention is all you need. In: Advances in neural information processing systems, pp 5999\u20136009. Neural information processing systems foundation"},{"key":"9498_CR41","doi-asserted-by":"crossref","unstructured":"Vedantam R, Zitnick CL, Parikh D (2014) CIDEr: consensus-based image description evaluation. In: Proceedings of the ieee computer society conference on computer vision and pattern recognition, pp 4566\u20134575","DOI":"10.1109\/CVPR.2015.7299087"},{"key":"9498_CR42","doi-asserted-by":"crossref","unstructured":"Veit A, Kovacs B, Bell S, McAuley J, Bala K, Belongie S (2015) Learning visual clothing style with heterogeneous dyadic co-occurrences. In: Proceedings of the IEEE international conference on computer vision, pp 4642\u20134650","DOI":"10.1109\/ICCV.2015.527"},{"key":"9498_CR43","doi-asserted-by":"publisher","unstructured":"Wang X, Wu B, Ye Y, Zhong Y (2019) Outfit compatibility prediction and diagnosis with multi-layered comparison network. In: MM 2019 \u2014Proceedings of the 27th ACM international conference on multimedia, pp 329\u2013337. https:\/\/doi.org\/10.1145\/3343031.3350909","DOI":"10.1145\/3343031.3350909"},{"key":"9498_CR44","unstructured":"Xu K, Ba JL, Kiros R, Cho K, Courville A, Salakhutdinov R, Zemel RS, Bengio Y (2015) Show, attend and tell: Neural image caption generation with visual attention. In: 32nd international conference on machine learning, ICML, pp 2048\u20132057. International Machine Learning Society (IMLS)"},{"key":"9498_CR46","doi-asserted-by":"crossref","unstructured":"Yang X, Zhang H, Jin D, Liu Y, Wu C-H, Tan J, Xie D, Wang J, Wang X (2020) Fashion captioning: towards generating accurate descriptions with semantic rewards. In: Lecture notes in computer science (including subseries lecture notes in artificial intelligence and lecture notes in bioinformatics). 12358 LNCS, pp 1\u201317","DOI":"10.1007\/978-3-030-58601-0_1"},{"key":"9498_CR45","doi-asserted-by":"publisher","first-page":"361","DOI":"10.1145\/3425636","volume":"17","author":"X Yang","year":"2021","unstructured":"Yang X, Song X, Feng F, Wen H, Duan L-Y, Nie L (2021) Attribute-wise Explainable Fashion Compatibility Modeling. ACM Trans Multimedia Comput Commun Appl 17:361\u20133621. https:\/\/doi.org\/10.1145\/3425636","journal-title":"ACM Trans Multimedia Comput Commun Appl"},{"key":"9498_CR47","doi-asserted-by":"publisher","first-page":"4519","DOI":"10.1007\/s00521-018-3691-y","volume":"32","author":"H Zhang","year":"2020","unstructured":"Zhang H, Sun Y, Liu L, Wang X, Li L, Liu W (2020) ClothingOut: a category-supervised GAN model for clothing segmentation and retrieval. Neural Comput Appl 32:4519\u20134530. https:\/\/doi.org\/10.1007\/s00521-018-3691-y","journal-title":"Neural Comput Appl"},{"key":"9498_CR48","doi-asserted-by":"publisher","unstructured":"Zheng S, Yang F, Kiapour M, Piramuthu R (2018) ModaNet: a large-scale street fashion dataset with polygon annotations. In: Presented at the October 15. https:\/\/doi.org\/10.1145\/3240508.3240652","DOI":"10.1145\/3240508.3240652"}],"container-title":["Evolving Systems"],"original-title":[],"language":"en","link":[{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s12530-023-09498-w.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/article\/10.1007\/s12530-023-09498-w\/fulltext.html","content-type":"text\/html","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/link.springer.com\/content\/pdf\/10.1007\/s12530-023-09498-w.pdf","content-type":"application\/pdf","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2024,6,4]],"date-time":"2024-06-04T11:18:21Z","timestamp":1717499901000},"score":1,"resource":{"primary":{"URL":"https:\/\/link.springer.com\/10.1007\/s12530-023-09498-w"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2023,4,8]]},"references-count":48,"journal-issue":{"issue":"3","published-print":{"date-parts":[[2024,6]]}},"alternative-id":["9498"],"URL":"https:\/\/doi.org\/10.1007\/s12530-023-09498-w","relation":{},"ISSN":["1868-6478","1868-6486"],"issn-type":[{"value":"1868-6478","type":"print"},{"value":"1868-6486","type":"electronic"}],"subject":[],"published":{"date-parts":[[2023,4,8]]},"assertion":[{"value":"18 December 2022","order":1,"name":"received","label":"Received","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"27 March 2023","order":2,"name":"accepted","label":"Accepted","group":{"name":"ArticleHistory","label":"Article History"}},{"value":"8 April 2023","order":3,"name":"first_online","label":"First Online","group":{"name":"ArticleHistory","label":"Article History"}},{"order":1,"name":"Ethics","group":{"name":"EthicsHeading","label":"Declarations"}},{"value":"The authors declare that they have no known competing financial interests or personal relationships that could have appeared to influence the work reported in this paper.","order":2,"name":"Ethics","group":{"name":"EthicsHeading","label":"Conflict of interests"}}]}}