{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2025,7,16]],"date-time":"2025-07-16T13:55:04Z","timestamp":1752674104825,"version":"3.41.0"},"publisher-location":"New York, NY, USA","reference-count":51,"publisher":"ACM","license":[{"start":{"date-parts":[[2016,2,18]],"date-time":"2016-02-18T00:00:00Z","timestamp":1455753600000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.acm.org\/publications\/policies\/copyright_policy#Background"}],"content-domain":{"domain":["dl.acm.org"],"crossmark-restriction":true},"short-container-title":[],"published-print":{"date-parts":[[2016,2,18]]},"DOI":"10.1145\/2856636.2856643","type":"proceedings-article","created":{"date-parts":[[2016,2,1]],"date-time":"2016-02-01T20:06:43Z","timestamp":1454357203000},"page":"37-47","update-policy":"https:\/\/doi.org\/10.1145\/crossmark-policy","source":"Crossref","is-referenced-by-count":3,"title":["Energy and Performance Prediction of CUDA Applications using Dynamic Regression Models"],"prefix":"10.1145","author":[{"given":"Shajulin","family":"Benedict","sequence":"first","affiliation":[{"name":"HPCCLoud Research Laboratory, SXCCE, Anna University, India"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"R. S.","family":"Rejitha","sequence":"additional","affiliation":[{"name":"HPCCLoud Research Laboratory, SXCCE, Anna University, India"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Suja A.","family":"Alex","sequence":"additional","affiliation":[{"name":"HPCCLoud Research Laboratory, SXCCE, Anna University, India"}],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"320","published-online":{"date-parts":[[2016,2,18]]},"reference":[{"key":"e_1_3_2_1_1_1","doi-asserted-by":"crossref","unstructured":"Alcides\n       \n      Fonseca Bruno\n       \n      Cabral \u00c3&Eogon;miniumGPU\n      \n  \n  : \n  An Intelligent Framework for GPU Programming Facing the Multicore-Challenge III Volume \n  7686\n   of \n  the series Lecture Notes in Computer Science pp \n  96\n  --\n  107 2013\n  .  Alcides Fonseca Bruno Cabral \u00c3&Eogon;miniumGPU: An Intelligent Framework for GPU Programming Facing the Multicore-Challenge III Volume 7686 of the series Lecture Notes in Computer Science pp 96--107 2013.","DOI":"10.1007\/978-3-642-35893-7_9"},{"key":"e_1_3_2_1_2_1","doi-asserted-by":"publisher","DOI":"10.1145\/1735688.1735696"},{"key":"e_1_3_2_1_3_1","doi-asserted-by":"crossref","unstructured":"Ariel A. Fung W.W.L. Turner A.E. Aamodt T.M. \"Visualizing complex dynamics in many-core accelerator architectures \" 2010 IEEE International Symposium on Performance Analysis of Systems & Software (ISPASS) pp.164--174 28--30 March 2010 doi: 10.1109\/ISPASS.2010.5452029.  Ariel A. Fung W.W.L. Turner A.E. Aamodt T.M. \"Visualizing complex dynamics in many-core accelerator architectures \" 2010 IEEE International Symposium on Performance Analysis of Systems & Software (ISPASS) pp.164--174 28--30 March 2010 doi: 10.1109\/ISPASS.2010.5452029.","DOI":"10.1109\/ISPASS.2010.5452029"},{"volume-title":"echReport.pdf?sequence=3&isAllowed=y, accessed in Dec.","year":"2015","author":"Aji Ashwin M.","key":"e_1_3_2_1_4_1"},{"key":"e_1_3_2_1_5_1","doi-asserted-by":"publisher","DOI":"10.1007\/s11227-005-2335-z"},{"key":"e_1_3_2_1_6_1","doi-asserted-by":"publisher","DOI":"10.1145\/1375527.1375580"},{"key":"e_1_3_2_1_7_1","unstructured":"Barra Simulator https:\/\/code.google.com\/p\/barra-sim\/ accessed on 1 Oct 2015.  Barra Simulator https:\/\/code.google.com\/p\/barra-sim\/ accessed on 1 Oct 2015."},{"key":"e_1_3_2_1_8_1","doi-asserted-by":"publisher","DOI":"10.1109\/IPDPSW.2013.236"},{"key":"e_1_3_2_1_9_1","unstructured":"AutoTuning Seminar http:\/\/www.dagstuhl.de\/en\/program\/calendar\/semhp\/?semnr=13401 accessed on Oct 2015.  AutoTuning Seminar http:\/\/www.dagstuhl.de\/en\/program\/calendar\/semhp\/?semnr=13401 accessed on Oct 2015."},{"key":"e_1_3_2_1_10_1","doi-asserted-by":"crossref","unstructured":"J. Brehm and P. Worley \"Performance Prediction for Complex Parallel Applications \" in 11th International Symposium on Parallel Processing 1997.   J. Brehm and P. Worley \"Performance Prediction for Complex Parallel Applications \" in 11th International Symposium on Parallel Processing 1997.","DOI":"10.2172\/467122"},{"volume-title":"USA","year":"2002","author":"Cao J.","key":"e_1_3_2_1_11_1"},{"key":"e_1_3_2_1_12_1","doi-asserted-by":"publisher","DOI":"10.1016\/j.future.2004.11.019"},{"key":"e_1_3_2_1_13_1","unstructured":"Cyril Zeller NVIDIA's vision for Exa-scale http:\/\/www.teratec.eu\/library\/pdf\/forum\/2014\/Presentations\/SP04_C_Zeller_NVIDIA_Forum_Teratec_2014.pdf.  Cyril Zeller NVIDIA's vision for Exa-scale http:\/\/www.teratec.eu\/library\/pdf\/forum\/2014\/Presentations\/SP04_C_Zeller_NVIDIA_Forum_Teratec_2014.pdf."},{"key":"e_1_3_2_1_14_1","doi-asserted-by":"publisher","DOI":"10.1145\/1984693.1984697"},{"key":"e_1_3_2_1_15_1","doi-asserted-by":"publisher","DOI":"10.1007\/s11227-015-1467-z"},{"key":"e_1_3_2_1_16_1","doi-asserted-by":"crossref","unstructured":"Fauzia N. Pouchet L.-N. Sadayappan P. \"Characterizing and enhancing global memory data coalescing on GPUs \" 2015 IEEE\/ACM International Symposium on Code Generation and Optimization (CGO) pp.12--22 7--11 Feb. 2015 doi: 10.1109\/CGO.2015.7054183.   Fauzia N. Pouchet L.-N. Sadayappan P. \"Characterizing and enhancing global memory data coalescing on GPUs \" 2015 IEEE\/ACM International Symposium on Code Generation and Optimization (CGO) pp.12--22 7--11 Feb. 2015 doi: 10.1109\/CGO.2015.7054183.","DOI":"10.1109\/CGO.2015.7054183"},{"key":"e_1_3_2_1_17_1","doi-asserted-by":"publisher","DOI":"10.1016\/j.parco.2015.02.002"},{"key":"e_1_3_2_1_18_1","doi-asserted-by":"crossref","unstructured":"Hartwig A. Blake H. Jakub K. Piotr L. and Dongarra J. Experiences in autotuning matrix multiplication for energy minimization on GPUs Concurrency and Computation: Practice and Experience DOI: 10.1002\/cpe.3516.    10.1002\/cpe.3516\nHartwig A. Blake H. Jakub K. Piotr L. and Dongarra J. Experiences in autotuning matrix multiplication for energy minimization on GPUs Concurrency and Computation: Practice and Experience DOI: 10.1002\/cpe.3516.","DOI":"10.1002\/cpe.3516"},{"key":"e_1_3_2_1_19_1","doi-asserted-by":"publisher","DOI":"10.1109\/12.817403"},{"key":"e_1_3_2_1_20_1","doi-asserted-by":"publisher","DOI":"10.1145\/2145816.2145819"},{"volume-title":"International Journal of Parallel Programming","author":"Jiayuan","key":"e_1_3_2_1_21_1"},{"key":"e_1_3_2_1_22_1","doi-asserted-by":"publisher","DOI":"10.1109\/CGO.2013.6494986"},{"key":"e_1_3_2_1_23_1","doi-asserted-by":"crossref","unstructured":"Kamil S. Chan C. Oliker L. Shalf J. and Williams S. \"An auto-tuning framework for parallel multicore stencil computations \" 2010 IEEE International Symposium on Parallel & Distributed Processing (IPDPS) pp.1--12 19--23 April 2010 doi: 10.1109\/IPDPS.2010.5470421.  Kamil S. Chan C. Oliker L. Shalf J. and Williams S. \"An auto-tuning framework for parallel multicore stencil computations \" 2010 IEEE International Symposium on Parallel & Distributed Processing (IPDPS) pp.1--12 19--23 April 2010 doi: 10.1109\/IPDPS.2010.5470421.","DOI":"10.1109\/IPDPS.2010.5470421"},{"key":"e_1_3_2_1_24_1","unstructured":"N. Kapadia J. Fortes and C. Brodley \"Predictive application-performance modeling in a computational grid environment \" in 8th International Symposium on High Performance Distributed Computing 1999.   N. Kapadia J. Fortes and C. Brodley \"Predictive application-performance modeling in a computational grid environment \" in 8th International Symposium on High Performance Distributed Computing 1999."},{"key":"e_1_3_2_1_25_1","unstructured":"H. Li D. Groep J. Templon and L. Wolters \"Predicting Job Start Times on Clusters \" in International Symposium on Cluster Computing and the Grid 2004.   H. Li D. Groep J. Templon and L. Wolters \"Predicting Job Start Times on Clusters \" in International Symposium on Cluster Computing and the Grid 2004."},{"key":"e_1_3_2_1_26_1","doi-asserted-by":"publisher","DOI":"10.1007\/978-3-642-11515-8_10"},{"key":"e_1_3_2_1_27_1","doi-asserted-by":"publisher","DOI":"10.1109\/E-SCIENCE.2006.131"},{"key":"e_1_3_2_1_28_1","unstructured":"NVIDIA GPGPU applications http:\/\/www.nvidia.com\/content\/gpu-applications\/PDF\/GPU-apps-catalog-mar2015.pdf accessed on 30 Sep. 2015.  NVIDIA GPGPU applications http:\/\/www.nvidia.com\/content\/gpu-applications\/PDF\/GPU-apps-catalog-mar2015.pdf accessed on 30 Sep. 2015."},{"volume-title":"7th IEEE IC3 2014","year":"2014","author":"Benedict Shajulin","key":"e_1_3_2_1_29_1"},{"key":"e_1_3_2_1_30_1","doi-asserted-by":"publisher","DOI":"10.1109\/IPDPSW.2015.12"},{"key":"e_1_3_2_1_31_1","unstructured":"Shajulin Benedict Rejitha R.S. Preethi B. Bright C. and Judyfer W.S. 'Energy Analysis of Code Regions of HPC Applications using EnergyAnalyzer Tool' (in press) Int. Journal of Computational Science and Engineering InderScience publishers 2015.  Shajulin Benedict Rejitha R.S. Preethi B. Bright C. and Judyfer W.S. 'Energy Analysis of Code Regions of HPC Applications using EnergyAnalyzer Tool' (in press) Int. Journal of Computational Science and Engineering InderScience publishers 2015."},{"key":"e_1_3_2_1_32_1","doi-asserted-by":"publisher","DOI":"10.1016\/j.jpdc.2008.05.014"},{"key":"e_1_3_2_1_33_1","doi-asserted-by":"publisher","DOI":"10.1109\/IPDPS.2013.73"},{"key":"e_1_3_2_1_34_1","doi-asserted-by":"publisher","DOI":"10.1109\/ICMLA.2010.52"},{"key":"e_1_3_2_1_35_1","doi-asserted-by":"publisher","DOI":"10.1016\/j.jpdc.2004.06.008"},{"key":"e_1_3_2_1_36_1","doi-asserted-by":"crossref","unstructured":"A. Snavely L. Carrington N. Wolter J. Labarta R. Badia and A. Purkayastha \"A framework for performance modeling and prediction \" in Supercomputing Conference 2002.   A. Snavely L. Carrington N. Wolter J. Labarta R. Badia and A. Purkayastha \"A framework for performance modeling and prediction \" in Supercomputing Conference 2002.","DOI":"10.1109\/SC.2002.10004"},{"key":"e_1_3_2_1_37_1","doi-asserted-by":"publisher","DOI":"10.1007\/s10586-007-0039-2"},{"key":"e_1_3_2_1_38_1","unstructured":"R. Susukita H. Ando M. Aoyagi H. Honda Y. Inadomi K. Inoue S. Ishizuki Y. Kimura H. Komatsu M. Kurokawa K. J. Murakami H. Shibamura S. Yamamura and Y. Yu \"Performance prediction of large-scale parallell system and application using macro-level simulation \" in Supercomputing Conference 2008.   R. Susukita H. Ando M. Aoyagi H. Honda Y. Inadomi K. Inoue S. Ishizuki Y. Kimura H. Komatsu M. Kurokawa K. J. Murakami H. Shibamura S. Yamamura and Y. Yu \"Performance prediction of large-scale parallell system and application using macro-level simulation \" in Supercomputing Conference 2008."},{"key":"e_1_3_2_1_39_1","unstructured":"A. Takefusa O. Tatebe S. Matsuoka and Y. Morita \"Performance analysis of scheduling and replication algorithms on Grid Datafarm architecture for high-energy physics applications \" in 12th International Symposium on High Performance Distributed Computing 2003.   A. Takefusa O. Tatebe S. Matsuoka and Y. Morita \"Performance analysis of scheduling and replication algorithms on Grid Datafarm architecture for high-energy physics applications \" in 12th International Symposium on High Performance Distributed Computing 2003."},{"volume-title":"USA","year":"2002","author":"Taylor V.","key":"e_1_3_2_1_40_1"},{"volume-title":"Applications and Services","year":"2006","author":"Wu X.","key":"e_1_3_2_1_41_1"},{"key":"e_1_3_2_1_42_1","doi-asserted-by":"publisher","DOI":"10.1109\/IPDPS.2014.103"},{"key":"e_1_3_2_1_43_1","doi-asserted-by":"publisher","DOI":"10.1007\/11428831_66"},{"key":"e_1_3_2_1_44_1","first-page":"135","article-title":"B, Optimizing ELARS Algorithms using NVIDIA CUDA Heterogeneous Parallel Programming Platform","author":"Vedran M., Martina","year":"2014","journal-title":"ICT Innovations"},{"key":"e_1_3_2_1_45_1","doi-asserted-by":"crossref","unstructured":"F. Vraalsen R. A. Aydt C. L. Mendes and D. A. Reed \"Performance Contracts: Predicting and Monitoring Grid Application Behavior \" in 2nd International Workshop on Grid Computing 2001.   F. Vraalsen R. A. Aydt C. L. Mendes and D. A. Reed \"Performance Contracts: Predicting and Monitoring Grid Application Behavior \" in 2nd International Workshop on Grid Computing 2001.","DOI":"10.1007\/3-540-45644-9_15"},{"key":"e_1_3_2_1_46_1","doi-asserted-by":"publisher","DOI":"10.1145\/1375783.1375813"},{"key":"e_1_3_2_1_47_1","doi-asserted-by":"publisher","DOI":"10.1109\/ICPP.2011.53"},{"key":"e_1_3_2_1_48_1","doi-asserted-by":"publisher","DOI":"10.1109\/SC.2005.20"},{"key":"e_1_3_2_1_49_1","doi-asserted-by":"publisher","DOI":"10.1016\/j.peva.2005.01.008"},{"key":"e_1_3_2_1_50_1","doi-asserted-by":"publisher","DOI":"10.1145\/2024724.2024754"},{"key":"e_1_3_2_1_51_1","first-page":"547","volume-title":"ICCSA 2011","author":"Kubota Yuji","year":"2011"}],"event":{"name":"ISEC '16: 9th India Software Engineering Conference","sponsor":["iSOFT iSOFT","ACM India ACM India"],"location":"Goa India","acronym":"ISEC '16"},"container-title":["Proceedings of the 9th India Software Engineering Conference"],"original-title":[],"link":[{"URL":"https:\/\/dl.acm.org\/doi\/10.1145\/2856636.2856643","content-type":"unspecified","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/dl.acm.org\/doi\/pdf\/10.1145\/2856636.2856643","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,6,18]],"date-time":"2025-06-18T04:54:12Z","timestamp":1750222452000},"score":1,"resource":{"primary":{"URL":"https:\/\/dl.acm.org\/doi\/10.1145\/2856636.2856643"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2016,2,18]]},"references-count":51,"alternative-id":["10.1145\/2856636.2856643","10.1145\/2856636"],"URL":"https:\/\/doi.org\/10.1145\/2856636.2856643","relation":{},"subject":[],"published":{"date-parts":[[2016,2,18]]},"assertion":[{"value":"2016-02-18","order":2,"name":"published","label":"Published","group":{"name":"publication_history","label":"Publication History"}}]}}