@misc{indiciae95a59e4ad53a, title = {SMILE: Infusing Spatial and Motion Semantics in Masked Video Learning}, author = {Fida Mohammad Thoker and Letian Jiang and Chen Zhao and Bernard Ghanem}, year = {2025}, url = {https://arxiv.org/abs/2504.00527}, note = {Source identifier: 2504.00527} }