2024
Wei, Xiwen; Li, Guihong; Marculescu, Radu
Fairness Implications of Machine Unlearning: Bias Risks in Removing NSFW Content from Text-to-Image Models Workshop Forthcoming
Forthcoming.
Abstract | BibTeX | Tags: Machine Unlearning, Trustworthy ML
@workshop{FAIRNESS_NEURIPS_2024,
title = {Fairness Implications of Machine Unlearning: Bias Risks in Removing NSFW Content from Text-to-Image Models},
author = {Xiwen Wei and Guihong Li and Radu Marculescu},
year = {2024},
date = {2024-12-15},
abstract = {The rapid development of large-scale text-to-image generative models has raised significant concerns about their potential misuse in generating harmful, misleading, or inappropriate content. To address these safety issues, various machine unlearning methods have been proposed to efficiently remove not-safe-for-work (NSFW) content without the need for complete model re-training. While these unlearning methods effectively enhance model safety, their impact on model fairness remains largely unexplored. In this paper, we examine the fairness implications of NSFW content removal via machine unlearning and discover that some methods can unintentionally amplify existing biases, increasing them by up to 6x. Our findings reveal that this increased bias arises from the biased synthetic training data used during the unlearning process. To mitigate this bias, we employ Bayesian optimization to identify the optimal training data composition, thus balancing safety and fairness.},
keywords = {Machine Unlearning, Trustworthy ML},
pubstate = {forthcoming},
tppubtype = {workshop}
}
Li, Guihong; Hsu, Hsiang; Chen, Chun-Fu; Marculescu, Radu
Fast-NTK: Parameter-Efficient Unlearning for Large-Scale Models Proceedings
Conference on Computer Vision and Pattern Recognition Workshop, 2024.
Links | BibTeX | Tags: Machine Unlearning, Trustworthy ML
@proceedings{nokey,
title = {Fast-NTK: Parameter-Efficient Unlearning for Large-Scale Models},
author = {Guihong Li and Hsiang Hsu and Chun-Fu Chen and Radu Marculescu},
url = {https://arxiv.org/abs/2312.14923},
year = {2024},
date = {2024-06-17},
howpublished = {Conference on Computer Vision and Pattern Recognition Workshop},
keywords = {Machine Unlearning, Trustworthy ML},
pubstate = {published},
tppubtype = {proceedings}
}
Li, Guihong; Hsu, Hsiang; Chen, Chun-Fu; Marculescu, Radu
Machine Unlearning for Image-to-Image Generative Models Proceedings
International Conference on Learning Representations, 2024.
Links | BibTeX | Tags: Featured, Generative AI, Machine Unlearning, Trustworthy ML
@proceedings{machine_unlearn,
title = {Machine Unlearning for Image-to-Image Generative Models},
author = {Guihong Li and Hsiang Hsu and Chun-Fu Chen and Radu Marculescu},
url = {https://arxiv.org/abs/2402.00351},
year = {2024},
date = {2024-05-07},
urldate = {2024-05-07},
howpublished = {International Conference on Learning Representations},
keywords = {Featured, Generative AI, Machine Unlearning, Trustworthy ML},
pubstate = {published},
tppubtype = {proceedings}
}


