@misc{indiciae80ec7d55439c, title = {SAVE: Speech-Aware Video Representation Learning for Video-Text Retrieval}, author = {Ruixiang Zhao and Zhihao Xu and Bangxiang Lan and Zijie Xin and Jingyu Liu and Xirong Li}, year = {2026}, url = {https://arxiv.org/abs/2603.08224}, note = {Source identifier: 2603.08224} }