@misc{indiciae2bbb4739ba1f, title = {ViTA: Visual-Linguistic Translation by Aligning Object Tags}, author = {Kshitij Gupta and Devansh Gautam and Radhika Mamidi}, year = {2021}, url = {https://arxiv.org/abs/2106.00250}, note = {Source identifier: 2106.00250} }