@misc{indiciae5db310207fb2, title = {Spatio-Temporal Attention Models for Grounded Video Captioning}, author = {Mihai Zanfir and Elisabeta Marinoiu and Cristian Sminchisescu}, year = {2016}, url = {https://arxiv.org/abs/1610.04997}, note = {Source identifier: 1610.04997} }