@misc{indiciaef5f6a601725f, title = {MCoT-MVS: Multi-level Vision Selection by Multi-modal Chain-of-Thought Reasoning for Composed Image Retrieval}, author = {Xuri Ge and Chunhao Wang and Xindi Wang and Zheyun Qin and Zhumin Chen and Xin Xin}, year = {2026}, url = {https://arxiv.org/abs/2603.17360}, note = {Source identifier: 2603.17360} }