@inproceedings{6117f66506924a6ebc955581831b26d7,
title = "Accelerating Convergence in Bounding Box Regression with a Refined IoU Loss Function",
abstract = "Bounding box regression (BBR) is a critical component in object detection, significantly influencing the accuracy of object localization. However, existing Intersection over Union (IoU)-based loss functions encounter two primary challenges: (i) The penalty factor configuration often results in the expansion of anchor boxes during the regression, which in turn slows the convergence rate of the loss. (ii) There is a spatial imbalance caused by the disproportionate influence of anchor boxes with minimal overlap with the ground truth boxes. To resolve these two challenges, this paper proposes a novel loss function termed Fast-IoU, designed to swiftly and precisely measure the overlap area and aspect ratio in BBR. Building upon this, a dynamic non-monotonic focusing mechanism is integrated to evaluate the quality of anchor boxes in a non-linear manner. Fast-IoU can enhance the capability to focus on anchor boxes of medium quality. By incorporating Fast-IoU into popular object detectors such as YOLOv7, YOLOv8 and YOLOv10, we achieved an increase in average precision and improved performance compared to their original loss functions on the MS COCO datasets, thus validating the effectiveness of ourproposed improvement strategies.",
keywords = "Bounding box regression, Focusing mechanism, IoU loss, Object detection, Spatial imbalance",
author = "Enhui Chai and Xingyu Li and Tianxiang Cui and Zheng Lu and Tesema, {Fiseha Berhanu}",
note = "Publisher Copyright: {\textcopyright} 2025 IEEE.; 2025 IEEE International Conference on Acoustics, Speech, and Signal Processing, ICASSP 2025 ; Conference date: 06-04-2025 Through 11-04-2025",
year = "2025",
doi = "10.1109/ICASSP49660.2025.10889366",
language = "English",
series = "ICASSP, IEEE International Conference on Acoustics, Speech and Signal Processing - Proceedings",
publisher = "Institute of Electrical and Electronics Engineers Inc.",
editor = "Rao, {Bhaskar D} and Isabel Trancoso and Gaurav Sharma and Mehta, {Neelesh B.}",
booktitle = "2025 IEEE International Conference on Acoustics, Speech, and Signal Processing, ICASSP 2025 - Proceedings",
address = "United States",
}