@misc{indiciae46ceca393930, title = {Principles of Visual Tokens for Efficient Video Understanding}, author = {Xinyue Hao and Gen Li and Shreyank N Gowda and Robert B Fisher and Jonathan Huang and Anurag Arnab and Laura Sevilla-Lara}, year = {2025}, url = {https://arxiv.org/abs/2411.13626}, note = {Source identifier: 2411.13626} }