@misc{indiciae6170200cf4e8, title = {HVD: Human Vision-Driven Video Representation Learning for Text-Video Retrieval}, author = {Zequn Xie and Xin Liu and Boyun Zhang and Yuxiao Lin and Sihang Cai and Tao Jin}, year = {2026}, url = {https://arxiv.org/abs/2601.16155}, note = {Source identifier: 2601.16155} }