@misc{indiciaee37e87e2415e, title = {SF2T: Self-supervised Fragment Finetuning of Video-LLMs for Fine-Grained Understanding}, author = {Yangliu Hu and Zikai Song and Na Feng and Yawei Luo and Junqing Yu and Yi-Ping Phoebe Chen and Wei Yang}, year = {2025}, url = {https://arxiv.org/abs/2504.07745}, note = {Source identifier: 2504.07745} }