###################################################################### # @CCOSTAN - Follow Me on X # For more info visit https://www.vcloudinfo.com/click-here # Original Repo : https://github.com/CCOSTAN/Home-AssistantConfig # ------------------------------------------------------------------- # Proxmox Host Automations - reboots, repairs, and Joanna dispatch # Nightly Frigate host reboot plus update/runtime/disk health automations. # ------------------------------------------------------------------- # Related Issue: 1584 # Related Issue: 1798 # Notes: Creates HA repair issues when proxmox nodes report updates. # Notes: Proxmox update activity also writes HA persistent notifications. # Notes: Adds normalized runtime + disk health signals for dashboard/alerts. # Notes: Joanna dispatch handles Proxmox updates plus sustained runtime/disk degradations. # Notes: Normalized disk usage sensors expose state_class for long-term trend rollups. ###################################################################### template: - sensor: - name: "Proxmox1 Disk Used Percentage" unique_id: proxmox1_disk_used_percentage unit_of_measurement: "%" state_class: measurement icon: mdi:harddisk availability: >- {% set preferred = states('sensor.node_proxmox1_disk_used_percentage') %} {% set used = states('sensor.node_proxmox1_disk') %} {% set total = states('sensor.node_proxmox1_max_disk') %} {{ preferred not in ['unknown', 'unavailable', 'none', ''] or (used not in ['unknown', 'unavailable', 'none', ''] and total not in ['unknown', 'unavailable', 'none', ''] and (total | float(0)) > 0) }} state: >- {% set preferred = states('sensor.node_proxmox1_disk_used_percentage') %} {% if preferred not in ['unknown', 'unavailable', 'none', ''] %} {{ preferred | float(0) | round(1) }} {% else %} {% set used = states('sensor.node_proxmox1_disk') | float(0) %} {% set total = states('sensor.node_proxmox1_max_disk') | float(0) %} {% if total > 0 %} {{ ((used / total) * 100) | round(1) }} {% else %} 0 {% endif %} {% endif %} - name: "Proxmox02 Disk Used Percentage" unique_id: proxmox02_disk_used_percentage unit_of_measurement: "%" state_class: measurement icon: mdi:harddisk availability: >- {% set preferred = states('sensor.node_proxmox02_disk_used_percentage') %} {% set used = states('sensor.node_proxmox02_disk') %} {% set total = states('sensor.node_proxmox02_max_disk') %} {{ preferred not in ['unknown', 'unavailable', 'none', ''] or (used not in ['unknown', 'unavailable', 'none', ''] and total not in ['unknown', 'unavailable', 'none', ''] and (total | float(0)) > 0) }} state: >- {% set preferred = states('sensor.node_proxmox02_disk_used_percentage') %} {% if preferred not in ['unknown', 'unavailable', 'none', ''] %} {{ preferred | float(0) | round(1) }} {% else %} {% set used = states('sensor.node_proxmox02_disk') | float(0) %} {% set total = states('sensor.node_proxmox02_max_disk') | float(0) %} {% if total > 0 %} {{ ((used / total) * 100) | round(1) }} {% else %} 0 {% endif %} {% endif %} - binary_sensor: - name: "Proxmox1 Runtime Healthy" unique_id: proxmox1_runtime_healthy device_class: running state: >- {% set state_value = states('binary_sensor.node_proxmox1_status') %} {% if state_value in ['on', 'off'] %} {{ state_value == 'on' }} {% else %} {% set status = states('sensor.node_proxmox1_status') | lower %} {{ status in ['online', 'running', 'on'] }} {% endif %} - name: "Proxmox02 Runtime Healthy" unique_id: proxmox02_runtime_healthy device_class: running state: >- {% set state_value = states('binary_sensor.node_proxmox02_status') %} {% if state_value in ['on', 'off'] %} {{ state_value == 'on' }} {% else %} {% set status = states('sensor.node_proxmox02_status') | lower %} {{ status in ['online', 'running', 'on'] }} {% endif %} automation: - alias: "Nightly Frigate Server Reboot" id: nightly_frigate_reboot description: "Reboots the Frigate server every day at 5 AM" mode: single trigger: - platform: time at: "05:00:00" action: - service: button.press target: entity_id: button.qemu_docker2_101_reboot - service: script.send_to_logbook data: topic: "FRIGATE" message: "Frigate server rebooted at 5 AM." - alias: "Proxmox Updates Repair Issues" id: proxmox_updates_repair description: "Track repair issues when Proxmox hosts report updates." mode: restart trigger: - platform: state entity_id: binary_sensor.node_proxmox1_updates_packages to: "on" - platform: state entity_id: binary_sensor.node_proxmox1_updates_packages to: "off" - platform: state entity_id: binary_sensor.node_proxmox02_updates_packages to: "on" - platform: state entity_id: binary_sensor.node_proxmox02_updates_packages to: "off" variables: node_name: > {% if 'proxmox1' in trigger.entity_id %}Proxmox1{% else %}Proxmox02{% endif %} issue_id: > {% if 'proxmox1' in trigger.entity_id %} proxmox1_updates_available {% else %} proxmox02_updates_available {% endif %} action: - choose: - conditions: "{{ trigger.to_state.state == 'on' }}" sequence: - service: repairs.create data: issue_id: "{{ issue_id }}" severity: warning persistent: false title: "{{ node_name }} has updates available" description: > {{ trigger.entity_id }} is ON, indicating pending updates on {{ node_name }}. Apply updates in Proxmox, then reload this sensor to clear the issue. - service: persistent_notification.create data: notification_id: "{{ issue_id }}" title: "{{ node_name }} Proxmox updates available" message: > {{ node_name }} reports pending Proxmox updates via {{ trigger.entity_id }}. Joanna upgrade orchestration will run after the cluster update state settles. default: - service: repairs.remove continue_on_error: true data: issue_id: "{{ issue_id }}" - service: persistent_notification.create data: notification_id: "{{ issue_id }}" title: "{{ node_name }} Proxmox updates cleared" message: > {{ node_name }} update telemetry returned to {{ trigger.to_state.state }}. Repairs state has been cleared and Home Assistant now considers this host patched. - service: script.send_to_logbook data: topic: "PROXMOX" message: "{{ node_name }} has been Patched" - alias: "Proxmox Updates Joanna Dispatch" id: proxmox_updates_joanna_dispatch description: "Dispatch Joanna when Proxmox host updates are available, with kernel-refresh routing hints." mode: restart trigger: - platform: state entity_id: - binary_sensor.node_proxmox1_updates_packages - binary_sensor.node_proxmox02_updates_packages to: "on" condition: - condition: template value_template: >- {{ is_state('binary_sensor.node_proxmox1_updates_packages', 'on') or is_state('binary_sensor.node_proxmox02_updates_packages', 'on') }} action: - delay: "00:01:00" - condition: template value_template: >- {{ is_state('binary_sensor.node_proxmox1_updates_packages', 'on') or is_state('binary_sensor.node_proxmox02_updates_packages', 'on') }} - variables: proxmox1_updates_count: "{{ states('sensor.node_proxmox1_total_updates') | int(0) }}" proxmox02_updates_count: "{{ states('sensor.node_proxmox02_total_updates') | int(0) }}" proxmox1_updates_summary: >- {% set updates = state_attr('sensor.node_proxmox1_total_updates', 'updates_list') | default([], true) %} {% if updates is sequence and updates is not string and updates | count > 0 %} {{ updates | join('; ') }} {% else %} none reported {% endif %} proxmox02_updates_summary: >- {% set updates = state_attr('sensor.node_proxmox02_total_updates', 'updates_list') | default([], true) %} {% if updates is sequence and updates is not string and updates | count > 0 %} {{ updates | join('; ') }} {% else %} none reported {% endif %} kernel_update_packages: >- {% set ns = namespace(items=[]) %} {% for entity_id in ['sensor.node_proxmox1_total_updates', 'sensor.node_proxmox02_total_updates'] %} {% set updates = state_attr(entity_id, 'updates_list') | default([], true) %} {% if updates is sequence and updates is not string %} {% for item in updates %} {% set haystack = item | string | lower %} {% if (haystack | regex_findall('(proxmox-kernel|pve-kernel|linux-image|kernel)') | count) > 0 %} {% set ns.items = ns.items + [item] %} {% endif %} {% endfor %} {% endif %} {% endfor %} {{ ns.items | join('; ') if ns.items | count > 0 else 'none detected' }} structured_request: |- PROXMOX_HOST_UPGRADE_REQUEST tracking_issue=https://github.com/CCOSTAN/Home-AssistantConfig/issues/1798 triggered_entity={{ trigger.entity_id }} cluster_nodes=ProxMox1,ProxMox02 proxmox1_updates_sensor=binary_sensor.node_proxmox1_updates_packages proxmox1_updates_state={{ states('binary_sensor.node_proxmox1_updates_packages') }} proxmox1_updates_count={{ proxmox1_updates_count }} proxmox1_updates={{ proxmox1_updates_summary }} proxmox02_updates_sensor=binary_sensor.node_proxmox02_updates_packages proxmox02_updates_state={{ states('binary_sensor.node_proxmox02_updates_packages') }} proxmox02_updates_count={{ proxmox02_updates_count }} proxmox02_updates={{ proxmox02_updates_summary }} kernel_update_detected={{ (kernel_update_packages | trim) != 'none detected' }} kernel_update_packages={{ kernel_update_packages | trim }} required_policy=Inspect live Proxmox cluster health, storage, VM placement, and node status before changing anything. Patch hosts one node at a time and verify each node returns healthy before moving to the next. If kernel packages are present or a reboot is required after kernel updates, use the $kernel-refresh skill and follow its ordered node-cycling workflow. Verify the live placement of Frigate/Docker14 VM 101, stop it before rebooting the host that owns it, and start it only after the host cycle and VM migrations are complete. Report progress, failures, resume state, and final verification. - service: script.send_to_logbook data: topic: "PROXMOX" message: >- Proxmox updates detected on one or more hosts. Joanna upgrade orchestration requested. Kernel refresh: {{ 'yes' if (kernel_update_packages | trim) != 'none detected' else 'not detected from HA package list' }}. - service: persistent_notification.create data: notification_id: "proxmox_updates_joanna_dispatch" title: "Proxmox updates detected" message: |- Joanna upgrade orchestration requested for Proxmox host updates. ProxMox1: {{ proxmox1_updates_count }} updates - {{ proxmox1_updates_summary | trim }} ProxMox02: {{ proxmox02_updates_count }} updates - {{ proxmox02_updates_summary | trim }} Kernel refresh: {% if (kernel_update_packages | trim) != 'none detected' %}yes - {{ kernel_update_packages | trim }}{% else %}not detected from HA package list{% endif %} - service: script.joanna_dispatch data: trigger_context: "HA automation proxmox_updates_joanna_dispatch (Proxmox Updates Joanna Dispatch)" source: "home_assistant_automation.proxmox_updates_joanna_dispatch" summary: >- Proxmox host updates detected. Kernel refresh {{ 'required' if (kernel_update_packages | trim) != 'none detected' else 'not detected from HA package list' }}. entity_ids: - "binary_sensor.node_proxmox1_updates_packages" - "sensor.node_proxmox1_total_updates" - "binary_sensor.node_proxmox02_updates_packages" - "sensor.node_proxmox02_total_updates" - "binary_sensor.proxmox1_runtime_healthy" - "binary_sensor.proxmox02_runtime_healthy" - "binary_sensor.qemu_docker2_101_status" diagnostics: >- proxmox1_updates_count={{ proxmox1_updates_count }}, proxmox02_updates_count={{ proxmox02_updates_count }}, kernel_update_detected={{ (kernel_update_packages | trim) != 'none detected' }}, kernel_update_packages={{ kernel_update_packages | trim }} request: "{{ structured_request }}" - alias: "Proxmox Runtime Repair Issues" id: proxmox_runtime_repairs description: "Create and clear Repairs when Proxmox node runtime becomes unhealthy." mode: restart trigger: - platform: state entity_id: - binary_sensor.proxmox1_runtime_healthy - binary_sensor.proxmox02_runtime_healthy variables: node_name: >- {% if 'proxmox1' in trigger.entity_id %}Proxmox1{% else %}Proxmox02{% endif %} issue_id: >- {% if 'proxmox1' in trigger.entity_id %} proxmox1_runtime_unhealthy {% else %} proxmox02_runtime_unhealthy {% endif %} runtime_entity: >- {% if 'proxmox1' in trigger.entity_id %} binary_sensor.proxmox1_runtime_healthy {% else %} binary_sensor.proxmox02_runtime_healthy {% endif %} status_entity: >- {% if 'proxmox1' in trigger.entity_id %} {% if states('binary_sensor.node_proxmox1_status') not in ['unknown', 'unavailable', 'none', ''] %} binary_sensor.node_proxmox1_status {% else %} sensor.node_proxmox1_status {% endif %} {% else %} {% if states('binary_sensor.node_proxmox02_status') not in ['unknown', 'unavailable', 'none', ''] %} binary_sensor.node_proxmox02_status {% else %} sensor.node_proxmox02_status {% endif %} {% endif %} status_value: "{{ states(status_entity) }}" trigger_context: "HA automation proxmox_runtime_repairs (Proxmox Runtime Repair Issues)" action: - choose: - conditions: "{{ trigger.to_state.state == 'off' }}" sequence: - delay: "00:02:00" - condition: template value_template: "{{ is_state(trigger.entity_id, 'off') }}" - service: repairs.create data: issue_id: "{{ issue_id }}" severity: error persistent: true title: "{{ node_name }} runtime degraded" description: > {{ node_name }} has remained offline for over 2 minutes. Check node status in Proxmox and restore runtime. - service: script.joanna_dispatch data: trigger_context: "{{ trigger_context }}" source: "home_assistant_automation.proxmox_runtime_repairs" summary: "{{ node_name }} runtime has remained degraded for over 2 minutes" entity_ids: - "{{ runtime_entity }}" - "{{ status_entity }}" diagnostics: >- issue_id={{ issue_id }}, node_name={{ node_name }}, runtime_entity={{ runtime_entity }}, status_entity={{ status_entity }}, status_value={{ status_value }}, unhealthy_for=2m request: >- Investigate {{ node_name }} runtime degradation and restore node availability if possible. Check host status, cluster connectivity, storage reachability, and recent update activity first. Do not reboot the host unless explicitly requested. - service: script.send_to_logbook data: topic: "PROXMOX" message: >- {{ node_name }} runtime is degraded. Repair {{ issue_id }} opened and Joanna investigation requested. default: - service: repairs.remove continue_on_error: true data: issue_id: "{{ issue_id }}" - service: script.send_to_logbook data: topic: "PROXMOX" message: "{{ node_name }} runtime recovered." - alias: "Proxmox Disk Pressure Repair Issues" id: proxmox_disk_pressure_repairs description: "Create and clear Repairs when Proxmox node disk usage stays elevated." mode: restart trigger: - platform: numeric_state entity_id: - sensor.proxmox1_disk_used_percentage - sensor.proxmox02_disk_used_percentage above: 85 below: 92 for: "00:15:00" id: warning - platform: numeric_state entity_id: - sensor.proxmox1_disk_used_percentage - sensor.proxmox02_disk_used_percentage above: 92 id: critical - platform: state entity_id: - sensor.proxmox1_disk_used_percentage - sensor.proxmox02_disk_used_percentage id: band_change - platform: numeric_state entity_id: - sensor.proxmox1_disk_used_percentage - sensor.proxmox02_disk_used_percentage below: 85 id: recovered variables: node_name: >- {% if 'proxmox1' in trigger.entity_id %}Proxmox1{% else %}Proxmox02{% endif %} issue_id: >- {% if 'proxmox1' in trigger.entity_id %} proxmox1_disk_pressure {% else %} proxmox02_disk_pressure {% endif %} disk_entity: "{{ trigger.entity_id }}" raw_disk_entity: >- {% if 'proxmox1' in trigger.entity_id %} sensor.node_proxmox1_disk_used_percentage {% else %} sensor.node_proxmox02_disk_used_percentage {% endif %} disk_pct: "{{ states(disk_entity) | float(0) }}" previous_disk_pct: >- {% if trigger.from_state is not none and trigger.from_state.state not in ['unknown', 'unavailable', 'none', ''] %} {{ trigger.from_state.state | float(0) }} {% else %} 0 {% endif %} previous_band: >- {% if previous_disk_pct >= 92 %} critical {% elif previous_disk_pct >= 85 %} warning {% else %} normal {% endif %} action: - choose: - conditions: - condition: trigger id: critical sequence: - service: repairs.create data: issue_id: "{{ issue_id }}" severity: error persistent: true title: "{{ node_name }} disk pressure critical ({{ disk_pct | round(1) }}%)" description: > {{ node_name }} disk usage is critically high. Free disk space or expand storage allocation. - service: script.joanna_dispatch data: trigger_context: "HA automation proxmox_disk_pressure_repairs (Proxmox Disk Pressure Repair Issues - Critical)" source: "home_assistant_automation.proxmox_disk_pressure_repairs.critical" summary: "{{ node_name }} disk pressure is critical at {{ disk_pct | round(1) }}%" entity_ids: - "{{ disk_entity }}" - "{{ raw_disk_entity }}" diagnostics: >- issue_id={{ issue_id }}, node_name={{ node_name }}, disk_entity={{ disk_entity }}, raw_disk_entity={{ raw_disk_entity }}, disk_pct={{ disk_pct | round(1) }}, threshold=92 request: >- Investigate critical disk pressure on {{ node_name }} and recommend safe remediation. Check local storage usage, backups, logs, snapshots, and VM or container disk consumers first. Do not delete VM disks or reboot the host unless explicitly requested. - service: script.send_to_logbook data: topic: "PROXMOX" message: >- {{ node_name }} disk usage is critical at {{ disk_pct | round(1) }}%. Repair {{ issue_id }} opened and Joanna investigation requested. - conditions: - condition: trigger id: warning - condition: template value_template: "{{ previous_band != 'critical' }}" sequence: - service: repairs.create data: issue_id: "{{ issue_id }}" severity: warning persistent: true title: "{{ node_name }} disk pressure warning ({{ disk_pct | round(1) }}%)" description: > {{ node_name }} disk usage has stayed above 85% for 15 minutes. Plan cleanup before capacity reaches critical levels. - service: script.joanna_dispatch data: trigger_context: "HA automation proxmox_disk_pressure_repairs (Proxmox Disk Pressure Repair Issues - Warning)" source: "home_assistant_automation.proxmox_disk_pressure_repairs.warning" summary: "{{ node_name }} disk pressure warning at {{ disk_pct | round(1) }}%" entity_ids: - "{{ disk_entity }}" - "{{ raw_disk_entity }}" diagnostics: >- issue_id={{ issue_id }}, node_name={{ node_name }}, disk_entity={{ disk_entity }}, raw_disk_entity={{ raw_disk_entity }}, disk_pct={{ disk_pct | round(1) }}, threshold=85, sustained_for=15m request: >- Investigate elevated disk usage on {{ node_name }} and recommend safe cleanup actions before it becomes critical. Check local storage usage, backups, logs, snapshots, and VM or container disk consumers first. Do not delete VM disks or reboot the host unless explicitly requested. - service: script.send_to_logbook data: topic: "PROXMOX" message: >- {{ node_name }} disk usage warning at {{ disk_pct | round(1) }}%. Repair {{ issue_id }} opened and Joanna investigation requested. - conditions: - condition: trigger id: band_change - condition: template value_template: "{{ previous_band == 'critical' and disk_pct >= 85 and disk_pct < 92 }}" sequence: - service: repairs.create data: issue_id: "{{ issue_id }}" severity: warning persistent: true title: "{{ node_name }} disk pressure warning ({{ disk_pct | round(1) }}%)" description: > {{ node_name }} disk usage is elevated but no longer critical. Plan cleanup before capacity reaches critical levels again. - conditions: - condition: trigger id: recovered sequence: - service: repairs.remove continue_on_error: true data: issue_id: "{{ issue_id }}"