@misc{indiciae7dba2989373f, title = {VIMPAC: Video Pre-Training via Masked Token Prediction and Contrastive Learning}, author = {Hao Tan and Jie Lei and Thomas Wolf and Mohit Bansal}, year = {2021}, url = {https://arxiv.org/abs/2106.11250}, note = {Source identifier: 2106.11250} }