@misc{indiciae5ac2b907a4c6, title = {Unified World Models: Coupling Video and Action Diffusion for Pretraining on Large Robotic Datasets}, author = {Chuning Zhu and Raymond Yu and Siyuan Feng and Benjamin Burchfiel and Paarth Shah and Abhishek Gupta}, year = {2025}, url = {https://arxiv.org/abs/2504.02792}, note = {Source identifier: 2504.02792} }