@misc{indiciae8e3bba4196d5, title = {Iterative Answer Prediction with Pointer-Augmented Multimodal Transformers for TextVQA}, author = {Ronghang Hu and Amanpreet Singh and Trevor Darrell and Marcus Rohrbach}, year = {2020}, url = {https://arxiv.org/abs/1911.06258}, note = {Source identifier: 1911.06258} }