@misc{indiciae4eab68a1913d, title = {VideoTok4D: A 4D-Aware Video Tokenizer for Compact World Representation}, author = {Xinyi Chen and Hanxin Zhu and Xijun Wang and Xingrui Wang and Sen Liang and Xin Li and Zhibo Chen}, year = {2026}, url = {https://arxiv.org/abs/2609.12874}, note = {Source identifier: 2609.12874} }