@misc{indiciae60b6416c7300, title = {Everything is a Video: Unifying Modalities through Next-Frame Prediction}, author = {G. Thomas Hudson and Dean Slack and Thomas Winterbottom and Jamie Sterling and Chenghao Xiao and Junjie Shentu and Noura Al Moubayed}, year = {2025}, url = {https://arxiv.org/abs/2411.10503}, note = {Source identifier: 2411.10503} }