@misc{indiciae50900dd4918c, title = {LinVT: Empower Your Image-level Large Language Model to Understand Videos}, author = {Lishuai Gao and Yujie Zhong and Yingsen Zeng and Haoxian Tan and Dengjie Li and Zheng Zhao}, year = {2024}, url = {https://arxiv.org/abs/2412.05185}, note = {Source identifier: 2412.05185} }