@misc{indiciae024d3704d1d5, title = {VideoSAVi: Self-Aligned Video Language Models without Human Supervision}, author = {Yogesh Kulkarni and Pooyan Fazli}, year = {2025}, url = {https://arxiv.org/abs/2412.00624}, note = {Source identifier: 2412.00624} }