@misc{indiciae2673dda2653e, title = {HVM-1: Large-scale video models pretrained with nearly 5000 hours of human-like video data}, author = {A. Emin Orhan}, year = {2024}, url = {https://arxiv.org/abs/2407.18067}, note = {Source identifier: 2407.18067} }