@misc{indiciaecee1d39f9333, title = {BOLT: Boost Large Vision-Language Model Without Training for Long-form Video Understanding}, author = {Shuming Liu and Chen Zhao and Tianqi Xu and Bernard Ghanem}, year = {2025}, url = {https://arxiv.org/abs/2503.21483}, note = {Source identifier: 2503.21483} }