@misc{indiciae4a235cb6f559, title = {Dense Image Representation with Spatial Pyramid VLAD Coding of CNN for Locally Robust Captioning}, author = {Andrew Shin and Masataka Yamaguchi and Katsunori Ohnishi and Tatsuya Harada}, year = {2016}, url = {https://arxiv.org/abs/1603.09046}, note = {Source identifier: 1603.09046} }