{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2025,11,18]],"date-time":"2025-11-18T12:16:37Z","timestamp":1763468197874,"version":"3.41.0"},"publisher-location":"New York, NY, USA","reference-count":41,"publisher":"ACM","license":[{"start":{"date-parts":[[2014,6,23]],"date-time":"2014-06-23T00:00:00Z","timestamp":1403481600000},"content-version":"vor","delay-in-days":0,"URL":"https:\/\/www.acm.org\/publications\/policies\/copyright_policy#Background"}],"content-domain":{"domain":["dl.acm.org"],"crossmark-restriction":true},"short-container-title":[],"published-print":{"date-parts":[[2014,6,23]]},"DOI":"10.1145\/2600212.2600228","type":"proceedings-article","created":{"date-parts":[[2014,6,20]],"date-time":"2014-06-20T13:06:05Z","timestamp":1403269565000},"page":"153-164","update-policy":"https:\/\/doi.org\/10.1145\/crossmark-policy","source":"Crossref","is-referenced-by-count":30,"title":["Design and evaluation of the gemtc framework for GPU-enabled many-task computing"],"prefix":"10.1145","author":[{"given":"Scott J.","family":"Krieder","sequence":"first","affiliation":[{"name":"Illinois Institute of Technology, Chicago, IL, USA"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Justin M.","family":"Wozniak","sequence":"additional","affiliation":[{"name":"Argonne National Laboratory, Argonne, IL, USA"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Timothy","family":"Armstrong","sequence":"additional","affiliation":[{"name":"University of Chicago, Chicago, IL, USA"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Michael","family":"Wilde","sequence":"additional","affiliation":[{"name":"University of Chicago &amp; Argonne National Laboratory, Argonne, IL, USA"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Daniel S.","family":"Katz","sequence":"additional","affiliation":[{"name":"University of Chicago &amp; Argonne National Laboratory, Chicago, IL, USA"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Benjamin","family":"Grimmer","sequence":"additional","affiliation":[{"name":"Illinois Institute of Technology, Chicago, IL, USA"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Ian T.","family":"Foster","sequence":"additional","affiliation":[{"name":"University of Chicago &amp; Argonne National Laboratory, Argonne, IL, USA"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Ioan","family":"Raicu","sequence":"additional","affiliation":[{"name":"Illinois Institute of Technology, Chicago, IL, USA"}],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"320","published-online":{"date-parts":[[2014,6,23]]},"reference":[{"key":"e_1_3_2_1_1_1","first-page":"1","volume-title":"SC '08","author":"Raicu I.","year":"2008","unstructured":"I. Raicu , Z. Zhang , M. Wilde , I. Foster , P. Beckman , K. Iskra , and B. Clifford , \" Toward loosely coupled programming on petascale systems,\" in Proc. of 2008 ACM\/IEEE Conf. on Supercomputing, ser . SC '08 . Piscataway, NJ: IEEE Press , 2008 , pp. 22: 1 -- 22 :12. I. Raicu, Z. Zhang, M. Wilde, I. Foster, P. Beckman, K. Iskra, and B. Clifford, \"Toward loosely coupled programming on petascale systems,\" in Proc. of 2008 ACM\/IEEE Conf. on Supercomputing, ser. SC '08. Piscataway, NJ: IEEE Press, 2008, pp. 22:1--22:12."},{"key":"e_1_3_2_1_2_1","volume-title":"ProQuest","author":"Raicu I.","year":"2009","unstructured":"I. Raicu , Many-task computing : bridging the gap between high-throughput computing and high-performance computing . ProQuest , 2009 . I. Raicu, Many-task computing: bridging the gap between high-throughput computing and high-performance computing. ProQuest, 2009."},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_3_1","DOI":"10.1145\/1362622.1362680"},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_4_1","DOI":"10.1016\/j.parco.2011.05.005"},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_5_1","DOI":"10.1007\/s10723-013-9259-2"},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_6_1","DOI":"10.1109\/UCC.2011.25"},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_7_1","DOI":"10.1109\/GRID.2004.14"},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_8_1","DOI":"10.1007\/10968987_3"},{"key":"e_1_3_2_1_9_1","volume-title":"an open source platform for hpc system software research,\" in Edinburgh BG\/L System Software Workshop","author":"Desai N.","year":"2005","unstructured":"N. Desai , \"Cobalt : an open source platform for hpc system software research,\" in Edinburgh BG\/L System Software Workshop , 2005 . N. Desai, \"Cobalt: an open source platform for hpc system software research,\" in Edinburgh BG\/L System Software Workshop, 2005."},{"key":"e_1_3_2_1_10_1","first-page":"80","article-title":"Sub-block jobs","author":"IBM","year":"2013","unstructured":"IBM , \" Sub-block jobs ,\" in IBM System Blue Gene Solution: Blue Gene\/Q System Administration , 2013 , pp. 80 -- 81 , Sec. 6.3. IBM, \"Sub-block jobs,\" in IBM System Blue Gene Solution: Blue Gene\/Q System Administration, 2013, pp. 80--81, Sec. 6.3.","journal-title":"IBM System Blue Gene Solution: Blue Gene\/Q System Administration"},{"key":"e_1_3_2_1_11_1","first-page":"14","volume-title":"\\The case for tiny tasks in compute clusters,\" in Proc. of the 14th USENIX Conf. on Hot Topics in Operating Systems","author":"Ousterhout K.","year":"2013","unstructured":"K. Ousterhout , A. Panda , J. Rosen , S. Venkataraman , R. Xin , S. Ratnasamy , S. Shenker , and I. Stoica , \\The case for tiny tasks in compute clusters,\" in Proc. of the 14th USENIX Conf. on Hot Topics in Operating Systems . USENIX Association , 2013 , pp. 14 -- 14 . K. Ousterhout, A. Panda, J. Rosen, S. Venkataraman, R. Xin, S. Ratnasamy, S. Shenker, and I. Stoica, \\The case for tiny tasks in compute clusters,\" in Proc. of the 14th USENIX Conf. on Hot Topics in Operating Systems. USENIX Association, 2013, pp. 14--14."},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_12_1","DOI":"10.1002\/9780470558027.ch13"},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_13_1","DOI":"10.1007\/978-3-642-32820-6_85"},{"key":"e_1_3_2_1_14_1","volume-title":"Asynchronous Concurrent Execution","author":"NVIDIA Inc.","year":"2013","unstructured":"NVIDIA Inc. , \" CUDA C Programming Guide PG-02829-001 v5.5, Section 3.2.5 , Asynchronous Concurrent Execution ,\" 2013 . NVIDIA Inc., \"CUDA C Programming Guide PG-02829-001 v5.5, Section 3.2.5, Asynchronous Concurrent Execution,\" 2013."},{"key":"e_1_3_2_1_15_1","volume-title":"Dynamic Parallelism Execution","author":"NVIDIA Inc.","year":"2013","unstructured":"NVIDIA Inc. , \" CUDA C Programming Guide PG-02829-001 v5.5, Appendix C , Dynamic Parallelism Execution ,\" 2013 . NVIDIA Inc., \"CUDA C Programming Guide PG-02829-001 v5.5, Appendix C, Dynamic Parallelism Execution,\" 2013."},{"key":"e_1_3_2_1_16_1","volume-title":"Understanding the costs of many-task computing workloads on intel xeon phi coprocessors,\" in 2nd Greater Chicago Area System Research Workshop (GCASR)","author":"Johnson J.","year":"2013","unstructured":"J. Johnson , S. J. Krieder , B. Grimmer , J. M. Wozniak , M. Wilde , and I. Raicu , \" Understanding the costs of many-task computing workloads on intel xeon phi coprocessors,\" in 2nd Greater Chicago Area System Research Workshop (GCASR) , 2013 . J. Johnson, S. J. Krieder, B. Grimmer, J. M. Wozniak, M. Wilde, and I. Raicu, \"Understanding the costs of many-task computing workloads on intel xeon phi coprocessors,\" in 2nd Greater Chicago Area System Research Workshop (GCASR), 2013."},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_17_1","DOI":"10.1109\/SERVICES.2007.63"},{"key":"e_1_3_2_1_18_1","volume-title":"CCGrid","author":"Wozniak J. M.","year":"2013","unstructured":"J. M. Wozniak , T. G. Armstrong , M. Wilde , D. S. Katz , E. Lusk , and I. T. Foster , \" Swift\/T: Scalable data ow programming for many-task applications,\" in Proc . CCGrid , 2013 . J. M. Wozniak, T. G. Armstrong, M. Wilde, D. S. Katz, E. Lusk, and I. T. Foster, \"Swift\/T: Scalable data ow programming for many-task applications,\" in Proc. CCGrid, 2013."},{"unstructured":"T. G. Armstrong J. M. Wozniak M. Wilde and I. T. Foster \"Compiler optimization for data-driven task parallelism on distributed memory systems \" ANL\/MCS-P5030--1013.  T. G. Armstrong J. M. Wozniak M. Wilde and I. T. Foster \"Compiler optimization for data-driven task parallelism on distributed memory systems \" ANL\/MCS-P5030--1013.","key":"e_1_3_2_1_19_1"},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_20_1","DOI":"10.3233\/FI-2013-949"},{"unstructured":"NCSA \"Blue Waters User Portal \" 2014 https:\/\/bluewaters.ncsa.illinois.edu\/hardware-summary.  NCSA \"Blue Waters User Portal \" 2014 https:\/\/bluewaters.ncsa.illinois.edu\/hardware-summary.","key":"e_1_3_2_1_21_1"},{"unstructured":"J. Burkardt \"MD - molecular dynamics \" 2013 http:\/\/people.sc.fsu.edu\/~jburkardt\/cppsrc\/md\/md.html.  J. Burkardt \"MD - molecular dynamics \" 2013 http:\/\/people.sc.fsu.edu\/~jburkardt\/cppsrc\/md\/md.html.","key":"e_1_3_2_1_22_1"},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_23_1","DOI":"10.1002\/pro.767"},{"key":"e_1_3_2_1_24_1","series-title":"Lecture Notes in Computational Science and Engineering","doi-asserted-by":"crossref","first-page":"103","DOI":"10.1007\/3-540-31618-3_7","volume-title":"Biomolecular sampling: Algorithms,test molecules, and metrics,\" in New Algorithms for Macromolecular Simulation","author":"Hampton S. S.","year":"2006","unstructured":"S. S. Hampton , P. Brenner , A. Wenger , S. Chatterjee , and J. A. Izaguirre , \" Biomolecular sampling: Algorithms,test molecules, and metrics,\" in New Algorithms for Macromolecular Simulation , ser. Lecture Notes in Computational Science and Engineering , B. Leimkuhler, C. Chipot, R. Elber, A. Laaksonen, A. Mark, T. Schlick, C. Sch\u00c3ijtte, and R. Skeel, Eds. Springer-Verlag, New York , 2006 , vol. 49 , pp. 103 -- 121 . S. S. Hampton, P. Brenner, A. Wenger, S. Chatterjee, and J. A. Izaguirre, \"Biomolecular sampling: Algorithms,test molecules, and metrics,\" in New Algorithms for Macromolecular Simulation, ser. Lecture Notes in Computational Science and Engineering, B. Leimkuhler, C. Chipot, R. Elber, A. Laaksonen, A. Mark, T. Schlick, C. Sch\u00c3ijtte, and R. Skeel, Eds. Springer-Verlag, New York, 2006, vol. 49, pp. 103--121."},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_25_1","DOI":"10.1145\/1941553.1941590"},{"key":"e_1_3_2_1_26_1","volume-title":"Symp. on Parallel & Distributed Processing (IPDPS). IEEE","author":"Chen L.","year":"2010","unstructured":"L. Chen , O. Villa , S. Krishnamoorthy , and G. R. Gao , \" Dynamic load balancing on single-and multi-gpu systems,\" in IEEE Intl . Symp. on Parallel & Distributed Processing (IPDPS). IEEE , 2010 . L. Chen, O. Villa, S. Krishnamoorthy, and G. R. Gao, \"Dynamic load balancing on single-and multi-gpu systems,\" in IEEE Intl. Symp. on Parallel & Distributed Processing (IPDPS). IEEE, 2010."},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_27_1","DOI":"10.1109\/CLUSTER.2011.50"},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_28_1","DOI":"10.1145\/2043556.2043579"},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_29_1","DOI":"10.1145\/2517349.2522715"},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_30_1","DOI":"10.1145\/1996130.1996160"},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_31_1","DOI":"10.1145\/2287076.2287090"},{"key":"e_1_3_2_1_32_1","first-page":"3","volume-title":"USENIXATC'11","author":"Gupta V.","year":"2011","unstructured":"V. Gupta , K. Schwan , N. Tolia , V. Talwar , and P. Ranganathan , \" Pegasus: coordinated scheduling for virtualized accelerator-based systems,\" in Proc. of the 2011 USENIX Annual Technical Conf., ser . USENIXATC'11 . Berkeley, CA, USA: USENIX Association , 2011 , pp. 3 -- 3 . V. Gupta, K. Schwan, N. Tolia, V. Talwar, and P. Ranganathan, \"Pegasus: coordinated scheduling for virtualized accelerator-based systems,\" in Proc. of the 2011 USENIX Annual Technical Conf., ser. USENIXATC'11. Berkeley, CA, USA: USENIX Association, 2011, pp. 3--3."},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_33_1","DOI":"10.1002\/cpe.1631"},{"key":"e_1_3_2_1_34_1","volume-title":"Symp. on Cluster, Cloud and Grid Computing (CCGrid). IEEE","author":"Zhang C.","year":"2013","unstructured":"C. Zhang , G. Han , and C.-L. Wang , \"GPU-TLS : An efficient runtime for speculative loop parallelization on gpus,\" in 13th IEEE\/ACM Intl . Symp. on Cluster, Cloud and Grid Computing (CCGrid). IEEE , 2013 . C. Zhang, G. Han, and C.-L. Wang, \"GPU-TLS: An efficient runtime for speculative loop parallelization on gpus,\" in 13th IEEE\/ACM Intl. Symp. on Cluster, Cloud and Grid Computing (CCGrid). IEEE, 2013."},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_35_1","DOI":"10.1007\/978-3-642-36036-7_14"},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_36_1","DOI":"10.1145\/2493123.2462921"},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_37_1","DOI":"10.1504\/IJCSE.2013.052110"},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_38_1","DOI":"10.1109\/IPDPS.2012.23"},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_39_1","DOI":"10.1145\/2063384.2063402"},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_40_1","DOI":"10.1145\/2555243.2555258"},{"doi-asserted-by":"publisher","key":"e_1_3_2_1_41_1","DOI":"10.1145\/2493123.2462915"}],"event":{"sponsor":["SIGARCH ACM Special Interest Group on Computer Architecture"],"acronym":"HPDC'14","name":"HPDC'14: The 23rd International Symposium on High-Performance Parallel and Distributed Computing","location":"Vancouver BC Canada"},"container-title":["Proceedings of the 23rd international symposium on High-performance parallel and distributed computing"],"original-title":[],"link":[{"URL":"https:\/\/dl.acm.org\/doi\/10.1145\/2600212.2600228","content-type":"unspecified","content-version":"vor","intended-application":"text-mining"},{"URL":"https:\/\/dl.acm.org\/doi\/pdf\/10.1145\/2600212.2600228","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2025,6,18]],"date-time":"2025-06-18T08:10:26Z","timestamp":1750234226000},"score":1,"resource":{"primary":{"URL":"https:\/\/dl.acm.org\/doi\/10.1145\/2600212.2600228"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2014,6,23]]},"references-count":41,"alternative-id":["10.1145\/2600212.2600228","10.1145\/2600212"],"URL":"https:\/\/doi.org\/10.1145\/2600212.2600228","relation":{},"subject":[],"published":{"date-parts":[[2014,6,23]]},"assertion":[{"value":"2014-06-23","order":2,"name":"published","label":"Published","group":{"name":"publication_history","label":"Publication History"}}]}}