@misc{indiciae2b92939a433f, title = {Building an Open-Vocabulary Video CLIP Model with Better Architectures, Optimization and Data}, author = {Zuxuan Wu and Zejia Weng and Wujian Peng and Xitong Yang and Ang Li and Larry S. Davis and Yu-Gang Jiang}, year = {2023}, url = {https://arxiv.org/abs/2310.05010}, note = {Source identifier: 2310.05010} }