@misc{indiciae9093680ffca6, title = {VEDIT: Latent Prediction Architecture For Procedural Video Representation Learning}, author = {Han Lin and Tushar Nagarajan and Nicolas Ballas and Mido Assran and Mojtaba Komeili and Mohit Bansal and Koustuv Sinha}, year = {2024}, url = {https://arxiv.org/abs/2410.03478}, note = {Source identifier: 2410.03478} }