@misc{indiciaec37b7861b812, title = {Vamos: Versatile Action Models for Video Understanding}, author = {Shijie Wang and Qi Zhao and Minh Quan Do and Nakul Agarwal and Kwonjoon Lee and Chen Sun}, year = {2024}, url = {https://arxiv.org/abs/2311.13627}, note = {Source identifier: 2311.13627} }