@misc{indiciae30d8bf7b598f, title = {Seg4Diff: Unveiling Open-Vocabulary Segmentation in Text-to-Image Diffusion Transformers}, author = {Chaehyun Kim and Heeseong Shin and Eunbeen Hong and Heeji Yoon and Anurag Arnab and Paul Hongsuck Seo and Sunghwan Hong and Seungryong Kim}, year = {2025}, url = {https://arxiv.org/abs/2509.18096}, note = {Source identifier: 2509.18096} }