@misc{indiciae1e1135dda4ae, title = {Multilevel Language and Vision Integration for Text-to-Clip Retrieval}, author = {Huijuan Xu and Kun He and Bryan A. Plummer and Leonid Sigal and Stan Sclaroff and Kate Saenko}, year = {2018}, url = {https://arxiv.org/abs/1804.05113}, note = {Source identifier: 1804.05113} }