@misc{indiciaede7e3e768e5f, title = {Memory-Augmented Vision-Language Agents for Persistent and Semantically Consistent Object Captioning}, author = {Tommaso Galliena and Stefano Rosa and Tommaso Apicella and Pietro Morerio and Alessio Del Bue and Lorenzo Natale}, year = {2026}, url = {https://arxiv.org/abs/2603.24257}, note = {Source identifier: 2603.24257} }