@misc{indiciaebaa9b8435f2e, title = {CaVe-VLM-CoT: An Interpretable Vision-Language Model Framework}, author = {Sneha Rao and Shaina Raza and Dhanesh Ramachandram}, year = {2026}, url = {https://arxiv.org/abs/2606.18385}, note = {Source identifier: 2606.18385} }