@misc{indiciae9b53b1fb6585, title = {How Much 3D Do Video Foundation Models Encode?}, author = {Zixuan Huang and Xiang Li and Zhaoyang Lv and James M. Rehg}, year = {2025}, url = {https://arxiv.org/abs/2512.19949}, note = {Source identifier: 2512.19949} }