@misc{indiciae9eee5577c2b8, title = {Video Representation Learning with Joint-Embedding Predictive Architectures}, author = {Katrina Drozdov and Ravid Shwartz-Ziv and Yann LeCun}, year = {2024}, url = {https://arxiv.org/abs/2412.10925}, note = {Source identifier: 2412.10925} }