@misc{indiciae8bef0a1108ed, title = {FashionM3: Multimodal, Multitask, and Multiround Fashion Assistant based on Unified Vision-Language Model}, author = {Kaicheng Pang and Xingxing Zou and Waikeung Wong}, year = {2025}, url = {https://arxiv.org/abs/2504.17826}, note = {Source identifier: 2504.17826} }