From 853d774959ec4962f3cf116d55b021f58939cee2 Mon Sep 17 00:00:00 2001 From: Christian Pointner Date: Wed, 22 Feb 2023 15:48:51 +0100 Subject: prometheus: raise warning level for hwmon temps --- roles/monitoring/prometheus/server/defaults/main/rules_node.yml | 2 +- 1 file changed, 1 insertion(+), 1 deletion(-) (limited to 'roles/monitoring/prometheus/server/defaults') diff --git a/roles/monitoring/prometheus/server/defaults/main/rules_node.yml b/roles/monitoring/prometheus/server/defaults/main/rules_node.yml index cb519cb6..d1368333 100644 --- a/roles/monitoring/prometheus/server/defaults/main/rules_node.yml +++ b/roles/monitoring/prometheus/server/defaults/main/rules_node.yml @@ -128,7 +128,7 @@ prometheus_server_rules_node: description: "The systemd service unit {{ '{{' }} $labels.name {{ '}}' }} is in failed state.\n VALUE = {{ '{{' }} $value {{ '}}' }}\n LABELS = {{ '{{' }} $labels {{ '}}' }}" - alert: HostPhysicalComponentTooHot - expr: node_hwmon_temp_celsius > 75 + expr: node_hwmon_temp_celsius > 85 for: 5m labels: severity: warning -- cgit v1.2.3