@misc{indiciae208521fb7135, title = {Scaling Video-Language Models to 10K Frames via Hierarchical Differential Distillation}, author = {Chuanqi Cheng and Jian Guan and Wei Wu and Rui Yan}, year = {2025}, url = {https://arxiv.org/abs/2504.02438}, note = {Source identifier: 2504.02438} }