% Publications of Carl Vondrick, Columbia University. % 99 entries, newest first. @article{rammohan2026lookthere, title={LookThere! Sparse Vision by Reinforced Selection}, author={Sreehari Rammohan and Yousef Yassin and Anthony Fuller and Junfeng Wen and Carl Vondrick and Evan Shelhamer}, journal={arXiv 2026}, year={2026}, url={https://arxiv.org/abs/2609.04698} } @article{sudhakar2026robot, title={Robot Critics that Sweat the Small Stuff}, author={Sruthi Sudhakar and Junbang Liang and Sreehari Rammohan and Pavel Tokmakov and Richard Zemel and Carl Vondrick}, journal={arXiv 2026}, year={2026}, url={https://robocritic.cs.columbia.edu} } @article{rammohan2026a, title={A²: Smaller Self-Supervised ViTs Localize Better than Larger Ones}, author={Sreehari Rammohan and Huy Ha and Carl Vondrick}, journal={arXiv 2026}, year={2026}, url={https://arxiv.org/abs/2606.03148} } @article{ramakrishnan2026do, title={Do multimodal models imagine electric sheep?}, author={Santhosh Kumar Ramakrishnan and Carl Vondrick and Raja Giryes and Philipp Krähenbühl and Vladlen Koltun}, journal={arXiv 2026}, year={2026}, url={https://arxiv.org/abs/2605.09693} } @inproceedings{mani2026fewshot, title={Few-Shot Design Optimization by Exploiting Auxiliary Information}, author={Arjun Mani and Carl Vondrick and Richard Zemel}, booktitle={ICML 2026}, year={2026}, url={https://designopt.cs.columbia.edu} } @article{ozguroglu2025new, title={New York Smells: A Large Multimodal Dataset for Olfaction}, author={Ege Ozguroglu and Junbang Liang and Ruoshi Liu and Mia Chiquier and Michael DeTienne and Wesley Wei Qian and Alexandra Horowitz and Andrew Owens and Carl Vondrick}, journal={arXiv 2025}, year={2025}, url={https://smell.cs.columbia.edu} } @article{liang2025video, title={Video Generators are Robot Policies}, author={Junbang Liang and Pavel Tokmakov and Ruoshi Liu and Sruthi Sudhakar and Paarth Shah and Rares Ambrus and Carl Vondrick}, journal={arXiv 2025}, year={2025}, url={https://arxiv.org/abs/2508.00795} } @inproceedings{nagrani2025minerva, title={MINERVA: Evaluating Complex Video Reasoning}, author={Arsha Nagrani and Sachit Menon and Ahmet Iscen and Shyamal Buch and Ramin Mehran and Nilpa Jha and Anja Hauth and Yukun Zhu and Carl Vondrick and Mikhail Sirotenko and Cordelia Schmid and Tobias Weyand}, booktitle={ICCV 2025}, year={2025}, url={https://arxiv.org/abs/2505.00681} } @article{kao2025towards, title={Towards LLM Agents for Earth Observation}, author={Chia Hsiang Kao and Wenting Zhao and Shreelekha Revankar and Samuel Speas and Snehal Bhagat and Rajeev Datta and Cheng Perng Phoo and Utkarsh Mall and Carl Vondrick and Kavita Bala and Bharath Hariharan}, journal={arXiv 2025}, year={2025}, url={https://arxiv.org/abs/2504.12110} } @article{chiquier2025teaching, title={Teaching Humans Subtle Differences with DIFF-usion}, author={Mia Chiquier and Orr Avrech and Yossi Gandelsman and Berthy Feng and Katherine Bouman and Carl Vondrick}, journal={arXiv 2025}, year={2025}, url={https://diff-usion.cs.columbia.edu} } @inproceedings{hayden2025generative, title={Generative Data Mining with Longtail-Guided Diffusion}, author={David S. Hayden and Mao Ye and Timur Garipov and Gregory P. Meyer and Carl Vondrick and Zhao Chen and Yuning Chai and Eric Wolff and Siddhartha S. Srinivasa}, booktitle={ICML 2025}, year={2025}, url={https://arxiv.org/abs/2502.01980} } @inproceedings{mall2025disciple, title={DiSciPLE: Learning Interpretable Programs for Scientific Visual Discovery}, author={Utkarsh Mall and Cheng Perng Phoo and Mia Chiquier and Bharath Hariharan and Kavita Bala and Carl Vondrick}, booktitle={CVPR 2025}, year={2025}, url={https://disciple.cs.columbia.edu} } @article{liu2025selfimproving, title={Self-Improving Autonomous Underwater Manipulation}, author={Ruoshi Liu and Huy Ha and Mengxue Hou and Shuran Song and Carl Vondrick}, journal={ICRA 2025}, year={2025}, url={https://aquabot.cs.columbia.edu} } @inproceedings{liu2024differentiable, title={Differentiable Robot Rendering}, author={Ruoshi Liu and Alper Canberk and Shuran Song and Carl Vondrick}, booktitle={CoRL 2024 (Oral)}, year={2024}, url={https://drrobot.cs.columbia.edu} } @inproceedings{liang2024dreamitate, title={Dreamitate: Real-World Visuomotor Policy Learning via Video Generation}, author={Junbang Liang and Ruoshi Liu and Ege Ozguroglu and Sruthi Sudhakar and Achal Dave and Pavel Tokmakov and Shuran Song and Carl Vondrick}, booktitle={CoRL 2024}, year={2024}, url={https://dreamitate.cs.columbia.edu} } @inproceedings{menon2024whiteboardofthought, title={Whiteboard-of-Thought: Thinking Step-by-Step Across Modalities}, author={Sachit Menon and Richard Zemel and Carl Vondrick}, booktitle={EMNLP 2024}, year={2024}, url={https://whiteboard.cs.columbia.edu} } @inproceedings{canberk2024erasedraw, title={EraseDraw: Learning to Draw Step-by-Step via Erasing Objects from Images}, author={Alper Canberk and Maksym Bondarenko and Ege Ozguroglu and Ruoshi Liu and Carl Vondrick}, booktitle={ECCV 2024}, year={2024}, url={https://erasedraw.cs.columbia.edu} } @inproceedings{sarin2024how, title={How Video Meetings Change Your Expression}, author={Sumit Sarin and Utkarsh Mall and Purva Tendulkar and Carl Vondrick}, booktitle={ECCV 2024}, year={2024}, url={https://facet.cs.columbia.edu} } @inproceedings{sudhakar2024controlling, title={Controlling the World by Sleight of Hand}, author={Sruthi Sudhakar and Ruoshi Liu and Basile Van Hoorick and Carl Vondrick and and Richard Zemel}, booktitle={ECCV 2024 (Oral)}, year={2024}, url={https://arxiv.org/pdf/2408.07147} } @inproceedings{hoorick2024generative, title={Generative Camera Dolly: Extreme Monocular Dynamic Novel View Synthesis}, author={Basile Van Hoorick and Rundi Wu and Ege Ozguroglu and Kyle Sargent and Ruoshi Liu and Pavel Tokmakov and Achal Dave and Changxi Zheng and Carl Vondrick}, booktitle={ECCV 2024 (Oral)}, year={2024}, url={https://gcd.cs.columbia.edu} } @inproceedings{chiquier2024evolving, title={Evolving Interpretable Visual Classifiers with Large Language Models}, author={Mia Chiquier and Utkarsh Mall and Carl Vondrick}, booktitle={ECCV 2024}, year={2024}, url={https://llm-mutate.cs.columbia.edu} } @inproceedings{chen2024selfie, title={SelfIE: Self-Interpretation of Large Language Model Embeddings}, author={Haozhe Chen and Carl Vondrick and Chengzhi Mao}, booktitle={ICML 2024}, year={2024}, url={https://selfie.cs.columbia.edu} } @article{liu2024paperbot, title={PaperBot: Learning to Design Real-World Tools Using Paper}, author={Ruoshi Liu and Junbang Liang and Sruthi Sudhakar and Huy Ha and Cheng Chi and Shuran Song and Carl Vondrick}, journal={arXiv 2024}, year={2024}, url={https://paperbot.cs.columbia.edu} } @inproceedings{ozguroglu2024pixgestalt, title={pix2gestalt: Amodal Segmentation by Synthesizing Wholes}, author={Ege Ozguroglu and Ruoshi Liu and Dídac Surís and Dian Chen and Achal Dave and Pavel Tokmakov and Carl Vondrick}, booktitle={CVPR 2024}, year={2024}, url={https://gestalt.cs.columbia.edu} } @inproceedings{mao2024raidar, title={Raidar: geneRative AI Detection viA Rewriting}, author={Chengzhi Mao and Carl Vondrick and Hao Wang and Junfeng Yang}, booktitle={ICLR 2024}, year={2024}, url={https://arxiv.org/pdf/2401.12970.pdf} } @inproceedings{chen2024interpreting, title={Interpreting and Controlling Vision Foundation Models via Text Explanations}, author={Haozhe Chen and Junfeng Yang and Carl Vondrick and Chengzhi Mao}, booktitle={ICLR 2024}, year={2024}, url={https://arxiv.org/abs/2310.10591} } @inproceedings{wu2024sindm, title={Sin3DM: Learning a Diffusion Model from a Single 3D Textured Shape}, author={Rundi Wu and Ruoshi Liu and Carl Vondrick and Changxi Zheng}, booktitle={ICLR 2024}, year={2024}, url={https://sin3dm.github.io} } @inproceedings{mall2024remote, title={Remote Sensing Vision-Language Foundation Models without Annotations via Ground Remote Alignment}, author={Utkarsh Mall and Cheng Perng Phoo and Meilin Liu and Carl Vondrick and Bharath Hariharan and Kavita Bala}, booktitle={ICLR 2024}, year={2024}, url={https://arxiv.org/pdf/2312.06960.pdf} } @inproceedings{deitke2023objaversexl, title={Objaverse-XL: A Universe of 10M+ 3D Objects}, author={Matt Deitke and others}, booktitle={NeurIPS 2023}, year={2023}, url={https://objaverse.allenai.org/objaverse-xl-paper.pdf} } @inproceedings{surís2023vipergpt, title={ViperGPT: Visual Inference via Python Execution for Reasoning}, author={Dídac Surís and Sachit Menon and Carl Vondrick}, booktitle={ICCV 2023 (Oral)}, year={2023}, url={https://viper.cs.columbia.edu} } @inproceedings{liu2023zeroto, title={Zero-1-to-3: Zero-shot One Image to 3D Object}, author={Ruoshi Liu and Rundi Wu and Basile Van Hoorick and Pavel Tokmakov and Sergey Zakharov and Carl Vondrick}, booktitle={ICCV 2023 (Oral)}, year={2023}, url={https://zero123.cs.columbia.edu} } @inproceedings{chiquier2023muscles, title={Muscles in Action}, author={Mia Chiquier and Carl Vondrick}, booktitle={ICCV 2023}, year={2023}, url={https://arxiv.org/abs/2212.02978} } @inproceedings{mani2023surfsup, title={SurfsUp: Learning Fluid Simulation for Novel Surfaces}, author={Arjun Mani and Ishaan Preetam Chandratreya and Elliot Creager and Carl Vondrick and Richard Zemel}, booktitle={ICCV 2023}, year={2023}, url={https://surfsup.cs.columbia.edu} } @inproceedings{liu2023landscape, title={Landscape Learning for Neural Network Inversion}, author={Ruoshi Liu and Chengzhi Mao and Purva Tendulkar and Hao Wang and Carl Vondrick}, booktitle={ICCV 2023}, year={2023}, url={https://arxiv.org/abs/2206.09027} } @inproceedings{chen2023shiftd, title={SHIFT3D: Synthesizing Hard Inputs For Tricking 3D Detectors}, author={Hongge Chen and Zhao Chen and Greg Meyer and Dennis Park and Carl Vondrick and Ashish Shrivastava and Yuning Chai}, booktitle={ICCV 2023}, year={2023}, url={https://openaccess.thecvf.com/content/ICCV2023/papers/Chen_SHIFT3D_Synthesizing_Hard_Inputs_For_Tricking_3D_Detectors_ICCV_2023_paper.pdf} } @inproceedings{mao2023robust, title={Robust Perception through Equivariance}, author={Chengzhi Mao and Lingyu Zhang and Abhishek Joshi and Junfeng Yang and Hao Wang and Carl Vondrick}, booktitle={ICML 2023}, year={2023}, url={https://arxiv.org/abs/2212.06079} } @inproceedings{liu2023humans, title={Humans as Light Bulbs: 3D Human Reconstruction from Thermal Reflection}, author={Ruoshi Liu and Carl Vondrick}, booktitle={CVPR 2023}, year={2023}, url={https://thermal.cs.columbia.edu} } @inproceedings{liu2023what, title={What You Can Reconstruct from a Shadow}, author={Ruoshi Liu and Sachit Menon and Chengzhi Mao and Dennis Park and Simon Stent and Carl Vondrick}, booktitle={CVPR 2023}, year={2023}, url={https://openaccess.thecvf.com//content/CVPR2023/papers/Liu_What_You_Can_Reconstruct_From_a_Shadow_CVPR_2023_paper.pdf} } @inproceedings{hoorick2023tracking, title={Tracking through Containers and Occluders in the Wild}, author={Basile Van Hoorick and Pavel Tokmakov and Simon Stent and Jie Li and Carl Vondrick}, booktitle={CVPR 2023}, year={2023}, url={https://tcow.cs.columbia.edu} } @inproceedings{tendulkar2023flex, title={FLEX: Full-Body Grasping Without Full-Body Grasps}, author={Purva Tendulkar and Dídac Surís and Carl Vondrick}, booktitle={CVPR 2023}, year={2023}, url={https://flex.cs.columbia.edu/} } @inproceedings{mao2023doubly, title={Doubly Right Object Recognition: A Why Prompt for Visual Rationales}, author={Chengzhi Mao and Revant Teotia and Amrutha Sundar and Sachit Menon and Junfeng Yang and Xin Wang and Carl Vondrick}, booktitle={CVPR 2023}, year={2023}, url={https://arxiv.org/abs/2212.06202} } @article{geng2023affective, title={Affective Faces for Goal-Driven Dyadic Communication}, author={Scott Geng and Revant Teotia and Purva Tendulkar and Sachit Menon and Carl Vondrick}, journal={arXiv 2023}, year={2023}, url={https://arxiv.org/abs/2301.10939} } @inproceedings{menon2023visual, title={Visual Classification via Description from Large Language Models}, author={Sachit Menon and Carl Vondrick}, booktitle={ICLR 2023 (Oral)}, year={2023}, url={https://arxiv.org/abs/2210.07183} } @inproceedings{mao2023understanding, title={Understanding Zero-Shot Adversarial Robustness for Large-Scale Models}, author={Chengzhi Mao and Scott Geng and Junfeng Yang and Xin Wang and Carl Vondrick}, booktitle={ICLR 2023}, year={2023}, url={https://arxiv.org/abs/2212.07016} } @article{zhang2022adversarially, title={Adversarially Robust Video Perception by Seeing Motion}, author={Lingyu Zhang and Chengzhi Mao and Junfeng Yang and Carl Vondrick}, journal={arXiv 2022}, year={2022}, url={https://arxiv.org/abs/2212.07815} } @article{menon2022task, title={Task Bias in Vision-Language Models}, author={Sachit Menon and Ishaan Preetam Chandratreya and Carl Vondrick}, journal={arXiv 2022}, year={2022}, url={https://arxiv.org/pdf/2212.04412.pdf} } @inproceedings{lu2022private, title={Private Multiparty Perception for Navigation}, author={Hui Lu and Mia Chiquier and Carl Vondrick}, booktitle={NeurIPS 2022}, year={2022}, url={https://arxiv.org/abs/2212.00912} } @inproceedings{surís2022representing, title={Representing Spatial Trajectories as Distributions}, author={Dídac Surís and Carl Vondrick}, booktitle={NeurIPS 2022}, year={2022}, url={https://trajectories.cs.columbia.edu} } @inproceedings{menon2022forgetmenot, title={Forget-me-not! Contrastive Critics for Mitigating Posterior Collapse}, author={Sachit Menon and David Blei and Carl Vondrick}, booktitle={UAI 2022}, year={2022}, url={https://openreview.net/pdf?id=SrgIkwLjql9} } @inproceedings{hoorick2022revealing, title={Revealing Occlusions with 4D Neural Fields}, author={Basile Van Hoorick and Purva Tendulkar and Dídac Surís and Dennis Park and Simon Stent and Carl Vondrick}, booktitle={CVPR 2022 (Oral)}, year={2022}, url={https://arxiv.org/pdf/2204.10916.pdf} } @inproceedings{surís2022globetrotter, title={Globetrotter: Connecting Languages by Connecting Images}, author={Dídac Surís and Dave Epstein and Carl Vondrick}, booktitle={CVPR 2022 (Oral)}, year={2022}, url={https://arxiv.org/pdf/2012.04631.pdf} } @inproceedings{mao2022causal, title={Causal Transportability for Visual Recognition}, author={Chengzhi Mao and Kevin Xia and James Wang and Hao Wang and Junfeng Yang and Elias Bareinboim and Carl Vondrick}, booktitle={CVPR 2022}, year={2022}, url={https://arxiv.org/abs/2204.12363} } @inproceedings{surís2022its, title={It's Time for Artistic Correspondence in Music and Video}, author={Dídac Surís and Carl Vondrick and Bryan Russell and Justin Salamon}, booktitle={CVPR 2022}, year={2022}, url={https://arxiv.org/pdf/2206.07148.pdf} } @inproceedings{price2022unweavenet, title={UnweaveNet: Unweaving Activity Stories}, author={Will Price and Carl Vondrick and Dima Damen}, booktitle={CVPR 2022}, year={2022}, url={https://arxiv.org/abs/2112.10194} } @inproceedings{fu2022there, title={There is a Time and Place for Reasoning Beyond the Image}, author={Xingyu Fu and Ben Zhou and Ishaan Preetam Chandratreya and Carl Vondrick and Dan Roth}, booktitle={ACL 2022 (Oral)}, year={2022}, url={https://arxiv.org/abs/2203.00758} } @inproceedings{chiquier2022realtime, title={Real-Time Neural Voice Camouflage}, author={Mia Chiquier and Chengzhi Mao and Carl Vondrick}, booktitle={ICLR 2022 (Oral)}, year={2022}, url={https://voicecamo.cs.columbia.edu} } @inproceedings{mao2022discrete, title={Discrete Representations Strengthen Vision Transformer Robustness}, author={Chengzhi Mao and Lu Jiang and Mostafa Dehghani and Carl Vondrick and Rahul Sukthankar and Irfan Essa}, booktitle={ICLR 2022}, year={2022}, url={https://arxiv.org/abs/2111.10493} } @article{chen2021fullbody, title={Full-Body Visual Self-Modeling of Robot Morphologies}, author={Boyuan Chen and Robert Kwiatkowski and Carl Vondrick and Hod Lipson}, journal={Science Robotics 2022}, year={2021}, url={https://robot-morphology.cs.columbia.edu} } @inproceedings{chen2021the, title={The Boombox: Visual Reconstruction from Acoustic Vibrations}, author={Boyuan Chen and Mia Chiquier and Hod Lipson and Carl Vondrick}, booktitle={CoRL 2021}, year={2021}, url={https://arxiv.org/abs/2105.08052} } @inproceedings{mao2021adversarial, title={Adversarial Attacks are Reversible with Natural Supervision}, author={Chengzhi Mao and Mia Chiquier and Hao Wang and Junfeng Yang and Carl Vondrick}, booktitle={ICCV 2021}, year={2021}, url={https://arxiv.org/abs/2103.14222} } @inproceedings{hoorick2021dissecting, title={Dissecting Image Crops}, author={Basile Van Hoorick and Carl Vondrick}, booktitle={ICCV 2021}, year={2021}, url={https://arxiv.org/pdf/2011.11831.pdf} } @inproceedings{surís2021learning, title={Learning the Predictability of the Future}, author={Dídac Surís and Ruoshi Liu and Carl Vondrick}, booktitle={CVPR 2021}, year={2021}, url={https://arxiv.org/pdf/2101.01600.pdf} } @inproceedings{mao2021generative, title={Generative Interventions for Causal Learning}, author={Chengzhi Mao and Amogh Gupta and Augustine Cha and Hao Wang and Junfeng Yang and Carl Vondrick}, booktitle={CVPR 2021}, year={2021}, url={https://arxiv.org/abs/2012.12265} } @inproceedings{epstein2021learning, title={Learning Goals from Failure}, author={Dave Epstein and Carl Vondrick}, booktitle={CVPR 2021}, year={2021}, url={https://arxiv.org/pdf/2006.15657.pdf} } @article{chen2021visual, title={Visual Behavior Modelling for Robotic Theory of Mind}, author={Boyuan Chen and Carl Vondrick and Hod Lipson}, journal={Scientific Reports 2021}, year={2021}, url={https://www.nature.com/articles/s41598-020-77918-x} } @inproceedings{xu2020listening, title={Listening to Sounds of Silence for Speech Denoising}, author={Ruilin Xu and Rundi Wu and Yuko Ishiwaka and Carl Vondrick and Changxi Zheng}, booktitle={NeurIPS 2020}, year={2020}, url={http://www.cs.columbia.edu/cg/listen_to_the_silence/paper.pdf} } @inproceedings{mao2020multitask, title={Multitask Learning Strengthens Adversarial Robustness}, author={Chengzhi Mao and Amogh Gupta and Vikram Nitin and Baishakhi Ray and Shuran Song and Junfeng Yang and Carl Vondrick}, booktitle={ECCV 2020 (Oral)}, year={2020}, url={https://arxiv.org/pdf/2007.07236.pdf} } @inproceedings{andonian2020we, title={We Have So Much In Common: Modeling Semantic Relational Set Abstractions in Videos}, author={Alex Andonian and Camilo Fosco and Mathew Monfort and Allen Lee and Carl Vondrick and Rogerio Feris}, booktitle={ECCV 2020}, year={2020}, url={https://arxiv.org/abs/2008.05596} } @inproceedings{surís2020learning, title={Learning to Learn Words from Visual Scenes}, author={Dídac Surís and Dave Epstein and Heng Ji and Shih-Fu Chang and Carl Vondrick}, booktitle={ECCV 2020}, year={2020}, url={https://arxiv.org/pdf/1911.11237.pdf} } @inproceedings{epstein2020oops, title={Oops! Predicting Unintentional Action in Video}, author={Dave Epstein and Boyuan Chen and Carl Vondrick}, booktitle={CVPR 2020}, year={2020}, url={https://arxiv.org/pdf/1911.11206.pdf} } @inproceedings{mao2019metric, title={Metric Learning for Adversarial Robustness}, author={Chengzhi Mao and Ziyuan Zhong and Junfeng Yang and Carl Vondrick and Baishakhi Ray}, booktitle={NeurIPS 2019}, year={2019}, url={https://arxiv.org/abs/1909.00900} } @inproceedings{sun2019videobert, title={VideoBERT: A Joint Model for Video and Language Representation Learning}, author={Chen Sun and Austin Myers and Carl Vondrick and Kevin Murphy and Cordelia Schmid}, booktitle={ICCV 2019}, year={2019}, url={https://arxiv.org/abs/1904.01766} } @inproceedings{akbari2019multilevel, title={Multi-level Multimodal Common Semantic Space for Image-Phrase Grounding}, author={Hassan Akbari and Svebor Karaman and Surabhi Bhargava and Brian Chen and Carl Vondrick and Shih-Fu Chang}, booktitle={CVPR 2019}, year={2019}, url={https://arxiv.org/pdf/1811.11683.pdf} } @inproceedings{sun2019relational, title={Relational Action Forecasting}, author={Chen Sun and Abhinav Shrivastava and Carl Vondrick and Rahul Sukthankar and Kevin Murphy and Cordelia Schmid}, booktitle={CVPR 2019}, year={2019}, url={https://arxiv.org/abs/1904.04231} } @article{monfort2019moments, title={Moments in Time Dataset: one million videos for event understanding}, author={Mathew Monfort and others}, journal={PAMI 2019}, year={2019}, url={http://moments.csail.mit.edu/} } @inproceedings{vondrick2018tracking, title={Tracking Emerges by Colorizing Videos}, author={Carl Vondrick and Abhinav Shrivastava and Alireza Fathi and Sergio Guadarrama and Kevin Murphy}, booktitle={ECCV 2018}, year={2018}, url={https://arxiv.org/pdf/1806.09594.pdf} } @inproceedings{zhao2018the, title={The Sound of Pixels}, author={Hang Zhao and Chuang Gan and Andrew Rouditchenko and Carl Vondrick and Josh McDermott and Antonio Torralba}, booktitle={ECCV 2018}, year={2018}, url={https://arxiv.org/abs/1804.03160} } @inproceedings{sun2018actorcentric, title={Actor-centric Relation Network}, author={Chen Sun and Abhinav Shrivastava and Carl Vondrick and Kevin Murphy and Rahul Sukthankar and Cordelia Schmid}, booktitle={ECCV 2018}, year={2018}, url={https://arxiv.org/pdf/1807.10982.pdf} } @inproceedings{gu2018ava, title={AVA: A Video Dataset of Spatio-temporally Localized Atomic Visual Actions}, author={Chunhui Gu and others}, booktitle={CVPR 2018 (Spotlight)}, year={2018}, url={https://arxiv.org/abs/1705.08421} } @inproceedings{recasens2017following, title={Following Gaze in Video}, author={Adria Recasens and Carl Vondrick and Aditya Khosla and Antonio Torralba}, booktitle={ICCV 2017}, year={2017}, url={https://www.cs.columbia.edu/~vondrick/videogaze.pdf} } @inproceedings{vondrick2017generating, title={Generating the Future with Adversarial Transformers}, author={Carl Vondrick and Antonio Torralba}, booktitle={CVPR 2017}, year={2017}, url={https://www.cs.columbia.edu/~vondrick/transformer.pdf} } @article{aytar2017crossmodal, title={Cross-Modal Scene Networks}, author={Yusuf Aytar and Lluis Castrejon and Carl Vondrick and Hamed Pirsiavash and Antonio Torralba}, journal={PAMI 2017}, year={2017}, url={https://www.cs.columbia.edu/~vondrick/cmplaces_pami.pdf} } @article{aytar2017see, title={See, Hear, and Read: Deep Aligned Representations}, author={Yusuf Aytar and Carl Vondrick and Antonio Torralba}, journal={arXiv 2017}, year={2017}, url={https://www.cs.columbia.edu/~vondrick/see-hear-read/} } @inproceedings{vondrick2016generating, title={Generating Videos with Scene Dynamics}, author={Carl Vondrick and Hamed Pirsiavash and Antonio Torralba}, booktitle={NeurIPS 2016}, year={2016}, url={https://www.cs.columbia.edu/~vondrick/tinyvideo/} } @inproceedings{aytar2016soundnet, title={SoundNet: Learning Sound Representations from Unlabeled Video}, author={Yusuf Aytar and Carl Vondrick and Antonio Torralba}, booktitle={NeurIPS 2016}, year={2016}, url={http://projects.csail.mit.edu/soundnet/} } @inproceedings{vondrick2016anticipating, title={Anticipating Visual Representations from Unlabeled Video}, author={Carl Vondrick and Hamed Pirsiavash and Antonio Torralba}, booktitle={CVPR 2016 (Spotlight)}, year={2016}, url={https://www.cs.columbia.edu/~vondrick/prediction} } @inproceedings{vondrick2016predicting, title={Predicting Motivations of Actions by Leveraging Text}, author={Carl Vondrick and Deniz Oktay and Hamed Pirsiavash and Antonio Torralba}, booktitle={CVPR 2016}, year={2016}, url={https://www.cs.columbia.edu/~vondrick/intention.pdf} } @inproceedings{castrejon2016learning, title={Learning Aligned Cross-Modal Representations from Weakly Aligned Data}, author={Lluis Castrejon and Yusuf Aytar and Carl Vondrick and Hamed Pirsiavash and Antonio Torralba}, booktitle={CVPR 2016}, year={2016}, url={https://www.cs.columbia.edu/~vondrick/adaptation.pdf} } @article{vondrick2016visualizing, title={Visualizing Object Detection Features}, author={Carl Vondrick and Aditya Khosla and Hamed Pirsiavash and Tomasz Malisiewicz and Antonio Torralba}, journal={IJCV 2016}, year={2016}, url={https://www.cs.columbia.edu/~vondrick/ihog} } @article{zhu2015do, title={Do We Need More Training Data?}, author={Xiangxin Zhu and Carl Vondrick and Charless C. Fowlkes and Deva Ramanan}, journal={IJCV 2015}, year={2015}, url={https://www.cs.columbia.edu/~vondrick/bigdata.pdf} } @inproceedings{vondrick2015learning, title={Learning Visual Biases from Human Imagination}, author={Carl Vondrick and Hamed Pirsiavash and Aude Oliva and Antonio Torralba}, booktitle={NeurIPS 2015}, year={2015}, url={https://www.cs.columbia.edu/~vondrick/imagination/paper.pdf} } @inproceedings{recasens2015where, title={Where are they looking?}, author={Adria Recasens and Aditya Khosla and Carl Vondrick and Antonio Torralba}, booktitle={NeurIPS 2015}, year={2015}, url={https://www.cs.columbia.edu/~vondrick/gaze.pdf} } @inproceedings{pirsiavash2014assessing, title={Assessing the Quality of Actions}, author={Hamed Pirsiavash and Carl Vondrick and Antonio Torralba}, booktitle={ECCV 2014}, year={2014}, url={https://www.cs.columbia.edu/~vondrick/quality.pdf} } @inproceedings{vondrick2013hoggles, title={HOGgles: Visualizing Object Detection Features}, author={Carl Vondrick and Aditya Khosla and Tomasz Malisiewicz and Antonio Torralba}, booktitle={ICCV 2013 (Oral)}, year={2013}, url={https://www.cs.columbia.edu/~vondrick/ihog} } @inproceedings{zhu2012do, title={Do We Need More Training Data or Better Models for Object Detection?}, author={Xiangxin Zhu and Carl Vondrick and Deva Ramanan and Charless C. Fowlkes}, booktitle={BMVC 2012}, year={2012}, url={https://www.cs.columbia.edu/~vondrick/largetrain.pdf} } @article{vondrick2012efficiently, title={Efficiently Scaling Up Crowdsourced Video Annotation}, author={Carl Vondrick and Donald Patterson and Deva Ramanan}, journal={IJCV 2012}, year={2012}, url={https://www.cs.columbia.edu/~vondrick/vatic/ijcv.pdf} } @inproceedings{vondrick2011video, title={Video Annotation and Tracking with Active Learning}, author={Carl Vondrick and Deva Ramanan}, booktitle={NeurIPS 2011}, year={2011}, url={https://www.cs.columbia.edu/~vondrick/vatic/videoalearn.pdf} } @inproceedings{oh2011a, title={A Large-scale Benchmark Dataset for Event Recognition}, author={Sangmin Oh and others}, booktitle={CVPR 2011}, year={2011}, url={https://www.cs.columbia.edu/~vondrick/vatic/virat.pdf} } @inproceedings{vondrick2010efficiently, title={Efficiently Scaling Up Video Annotation with Crowdsourced Marketplaces}, author={Carl Vondrick and Deva Ramanan and Donald Patterson}, booktitle={ECCV 2010}, year={2010}, url={https://www.cs.columbia.edu/~vondrick/vatic/scalingup.pdf} }