@misc{indiciae6af1e225017f, title = {VideoRAE: Taming Video Foundation Models for Generative Modeling via Representation Autoencoders}, author = {Zhihao Xie and Junfeng Wu and Xinting Hu and Junchao Huang and Li Jiang}, year = {2026}, url = {https://arxiv.org/abs/2607.14088}, note = {Source identifier: 2607.14088} }