@misc{indiciaefd90267d8dba, title = {Co-Grounding Networks with Semantic Attention for Referring Expression Comprehension in Videos}, author = {Sijie Song and Xudong Lin and Jiaying Liu and Zongming Guo and Shih-Fu Chang}, year = {2021}, url = {https://arxiv.org/abs/2103.12346}, note = {Source identifier: 2103.12346} }