@misc{indiciae9f0aa2e8ca73, title = {TVTSv2: Learning Out-of-the-box Spatiotemporal Visual Representations at Scale}, author = {Ziyun Zeng and Yixiao Ge and Zhan Tong and Xihui Liu and Shu-Tao Xia and Ying Shan}, year = {2023}, url = {https://arxiv.org/abs/2305.14173}, note = {Source identifier: 2305.14173} }