@misc{indiciae48fb577affe0, title = {CrossVideoMAE: Self-Supervised Image-Video Representation Learning with Masked Autoencoders}, author = {Shihab Aaqil Ahamed and Malitha Gunawardhana and Liel David and Michael Sidorov and Daniel Harari and Muhammad Haris Khan}, year = {2025}, url = {https://arxiv.org/abs/2502.07811}, note = {Source identifier: 2502.07811} }