@article{arxiv230518279,
  title     = {{Contextual Object Detection with Multimodal Large Language Models}},
  author    = {Yuhang Zang and Wei Li and Jun Han and Kaiyang Zhou and Chen Change Loy},
  journal   = {International Journal of Computer Vision},
  month     = {February},
  year      = {2025},
  volume    = {133},
  number    = {2},
  publisher = {Springer Science and Business Media LLC},
  pages     = {825--843},
  doi       = {10.1007/s11263-024-02214-4},
  url       = {https://doi.org/10.1007/s11263-024-02214-4}
}
