{"status":"ok","message-type":"work","message-version":"1.0.0","message":{"indexed":{"date-parts":[[2026,3,31]],"date-time":"2026-03-31T19:14:33Z","timestamp":1774984473923,"version":"3.50.1"},"reference-count":49,"publisher":"IEEE","license":[{"start":{"date-parts":[[2021,5,30]],"date-time":"2021-05-30T00:00:00Z","timestamp":1622332800000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-029"},{"start":{"date-parts":[[2021,5,30]],"date-time":"2021-05-30T00:00:00Z","timestamp":1622332800000},"content-version":"stm-asf","delay-in-days":0,"URL":"https:\/\/doi.org\/10.15223\/policy-037"}],"content-domain":{"domain":[],"crossmark-restriction":false},"short-container-title":[],"published-print":{"date-parts":[[2021,5,30]]},"DOI":"10.1109\/icra48506.2021.9562028","type":"proceedings-article","created":{"date-parts":[[2021,10,20]],"date-time":"2021-10-20T00:28:35Z","timestamp":1634689715000},"page":"13678-13685","source":"Crossref","is-referenced-by-count":4,"title":["Visual Perspective Taking for Opponent Behavior Modeling"],"prefix":"10.1109","author":[{"given":"Boyuan","family":"Chen","sequence":"first","affiliation":[{"name":"Columbia University"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Yuhang","family":"Hu","sequence":"additional","affiliation":[{"name":"Columbia University"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Robert","family":"Kwiatkowski","sequence":"additional","affiliation":[{"name":"Columbia University"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Shuran","family":"Song","sequence":"additional","affiliation":[{"name":"Columbia University"}],"role":[{"role":"author","vocabulary":"crossref"}]},{"given":"Hod","family":"Lipson","sequence":"additional","affiliation":[{"name":"Columbia University"}],"role":[{"role":"author","vocabulary":"crossref"}]}],"member":"263","reference":[{"key":"ref39","doi-asserted-by":"publisher","DOI":"10.1109\/IROS.2018.8593702"},{"key":"ref38","article-title":"Intentionnet: Integrating planning and deep learning for goal-directed autonomous navigation","author":"gao","year":"2017","journal-title":"arXiv preprint arXiv 1710 05627"},{"key":"ref33","article-title":"World models","author":"ha","year":"2018","journal-title":"arXiv preprint arXiv 1803 10122"},{"key":"ref32","article-title":"Dream to control: Learning behaviors by latent imagination","author":"hafner","year":"2019","journal-title":"arXiv preprint arXiv 1912 01603"},{"key":"ref31","doi-asserted-by":"publisher","DOI":"10.15607\/RSS.2020.XVI.034"},{"key":"ref30","doi-asserted-by":"publisher","DOI":"10.1109\/ICRA.2018.8461184"},{"key":"ref37","doi-asserted-by":"publisher","DOI":"10.1109\/TSSC.1968.300136"},{"key":"ref36","doi-asserted-by":"publisher","DOI":"10.1109\/ICRA.2011.5979561"},{"key":"ref35","doi-asserted-by":"publisher","DOI":"10.15607\/RSS.2020.XVI.035"},{"key":"ref34","doi-asserted-by":"publisher","DOI":"10.1109\/LRA.2019.2925731"},{"key":"ref28","article-title":"Emergent communication through negotiation","author":"cao","year":"2018","journal-title":"arXiv preprint arXiv 1804 03583"},{"key":"ref27","article-title":"Emergence of grounded compositional language in multi-agent populations","author":"mordatch","year":"2017","journal-title":"arXiv preprint arXiv 1703 04529"},{"key":"ref29","doi-asserted-by":"publisher","DOI":"10.1109\/CVPR.2019.00685"},{"key":"ref2","article-title":"The development of knowledge about visual perception","author":"flavell","year":"1977","journal-title":"Nebraska Symposium on Motivation"},{"key":"ref1","doi-asserted-by":"crossref","DOI":"10.4324\/9781315006239","volume":"4","author":"piaget","year":"2013","journal-title":"Child&#x2019;s Conception of Space Selected Works"},{"key":"ref20","doi-asserted-by":"publisher","DOI":"10.1126\/science.aax4705"},{"key":"ref22","first-page":"7","article-title":"A new approach to exploring language emergence as boundedly optimal control in the face of environmental and cognitive constraints","author":"bratman","year":"2010","journal-title":"Proceedings of the 10th International Conference on Cognitive Modeling"},{"key":"ref21","doi-asserted-by":"publisher","DOI":"10.1109\/TSMCA.2005.850592"},{"key":"ref24","article-title":"Multi-agent cooperation and the emergence of (natural) language","author":"lazaridou","year":"2016","journal-title":"arXiv preprint arXiv 1612 07182"},{"key":"ref23","first-page":"2137","article-title":"Learning to communicate with deep multi-agent reinforcement learning","author":"foerster","year":"2016","journal-title":"Advances in neural information processing systems"},{"key":"ref26","first-page":"2244","article-title":"Learning multiagent communication with backpropagation","author":"sukhbaatar","year":"2016","journal-title":"Advances in neural information processing systems"},{"key":"ref25","doi-asserted-by":"publisher","DOI":"10.1038\/srep34615"},{"key":"ref10","doi-asserted-by":"publisher","DOI":"10.1162\/isal_a_00269"},{"key":"ref11","article-title":"Emergent tool use from multi-agent autocurricula","author":"baker","year":"2020"},{"key":"ref40","article-title":"Follownet: Robot navigation by following natural language directions with deep reinforcement learning","author":"shah","year":"2018","journal-title":"arXiv preprint arXiv 1805 06150"},{"key":"ref12","article-title":"Artificial agents learn flexible visual representations by playing a hiding game","author":"weihs","year":"2019","journal-title":"arXiv preprint arXiv 1912 08195"},{"key":"ref13","doi-asserted-by":"publisher","DOI":"10.3389\/fncom.2020.00069"},{"key":"ref14","article-title":"A deep learning approach for joint video frame and reward prediction in atari games","author":"leibfried","year":"2016","journal-title":"arXiv preprint arXiv 1611 07078"},{"key":"ref15","doi-asserted-by":"publisher","DOI":"10.1109\/ICRA.2017.7989324"},{"key":"ref16","first-page":"5690","article-title":"Imaginationaugmented agents for deep reinforcement learning","author":"racaniere","year":"2017","journal-title":"Advances in neural information processing systems"},{"key":"ref17","article-title":"Recurrent environment simulators","author":"chiappa","year":"2017","journal-title":"arXiv preprint arXiv 1704 02254"},{"key":"ref18","first-page":"2555","article-title":"Learning latent dynamics for planning from pixels","author":"hafner","year":"2019","journal-title":"International Conference on Machine Learning"},{"key":"ref19","doi-asserted-by":"publisher","DOI":"10.3758\/s13414-017-1446-y"},{"key":"ref4","doi-asserted-by":"publisher","DOI":"10.1037\/0012-1649.28.4.635"},{"key":"ref3","doi-asserted-by":"publisher","DOI":"10.1037\/0012-1649.16.1.10"},{"key":"ref6","doi-asserted-by":"publisher","DOI":"10.3389\/fnhum.2013.00652"},{"key":"ref5","doi-asserted-by":"publisher","DOI":"10.1111\/j.1467-8624.2010.01571.x"},{"key":"ref8","doi-asserted-by":"publisher","DOI":"10.1145\/1121241.1121283"},{"key":"ref7","first-page":"1","article-title":"Visual behavior modelling for robotic theory of mind","volume":"11","author":"chen","year":"2021","journal-title":"Scientific Reports"},{"key":"ref49","doi-asserted-by":"crossref","first-page":"478","DOI":"10.1016\/j.procs.2018.07.060","article-title":"A survey on swarm robotic modeling, analysis and hardware architecture","volume":"133","author":"gs","year":"2018","journal-title":"Science Direct Procedia Computer Science"},{"key":"ref9","doi-asserted-by":"publisher","DOI":"10.1016\/j.artint.2018.01.002"},{"key":"ref46","article-title":"Unity: A general platform for intelligent agents","author":"juliani","year":"2018","journal-title":"arXiv preprint arXiv 1809 02627"},{"key":"ref45","article-title":"Visual foresight: Model-based deep reinforcement learning for vision-based robotic control","author":"ebert","year":"2018","journal-title":"arXiv preprint arXiv 1812 02588"},{"key":"ref48","doi-asserted-by":"publisher","DOI":"10.1109\/TRO.2018.2857475"},{"key":"ref47","article-title":"Learning and querying fast generative models for reinforcement learning","author":"buesing","year":"2018","journal-title":"arXiv preprint arXiv 1802 03395"},{"key":"ref42","first-page":"1113","article-title":"Learning latent plans from play","author":"lynch","year":"2020","journal-title":"Conference on Robot Learning"},{"key":"ref41","doi-asserted-by":"publisher","DOI":"10.1109\/IVS.2019.8814056"},{"key":"ref44","article-title":"Adam: A method for stochastic optimization","author":"kingma","year":"2014","journal-title":"arXiv preprint arXiv 1412 6980"},{"key":"ref43","doi-asserted-by":"publisher","DOI":"10.1109\/ICCV.2015.316"}],"event":{"name":"2021 IEEE International Conference on Robotics and Automation (ICRA)","location":"Xi'an, China","start":{"date-parts":[[2021,5,30]]},"end":{"date-parts":[[2021,6,5]]}},"container-title":["2021 IEEE International Conference on Robotics and Automation (ICRA)"],"original-title":[],"link":[{"URL":"http:\/\/xplorestaging.ieee.org\/ielx7\/9560720\/9560666\/09562028.pdf?arnumber=9562028","content-type":"unspecified","content-version":"vor","intended-application":"similarity-checking"}],"deposited":{"date-parts":[[2023,1,12]],"date-time":"2023-01-12T22:50:18Z","timestamp":1673563818000},"score":1,"resource":{"primary":{"URL":"https:\/\/ieeexplore.ieee.org\/document\/9562028\/"}},"subtitle":[],"short-title":[],"issued":{"date-parts":[[2021,5,30]]},"references-count":49,"URL":"https:\/\/doi.org\/10.1109\/icra48506.2021.9562028","relation":{},"subject":[],"published":{"date-parts":[[2021,5,30]]}}}