Trap handler behavior will differ when a debugger is attached.
Make the debug trap flag available in the trap handler TMA.
Update it when the debug trap ioctl is invoked.
Signed-off-by: Jay Cornwall <jay.cornwall@amd.com>
Reviewed-by: Felix Kuehling <Felix.Kuehling@amd.com>
Signed-off-by: Jonathan Kim <jonathan.kim@amd.com>
Signed-off-by: Alex Deucher <alexander.deucher@amd.com>
                if (unwind && i == unwind_count)
                        break;
 
+               kfd_process_set_trap_debug_flag(&pdd->qpd, false);
+
                /* GFX off is already disabled by debug activate if not RLC restore supported. */
                if (kfd_dbg_is_rlc_restore_supported(pdd->dev))
                        amdgpu_gfx_off_ctrl(pdd->dev->adev, false);
                if (kfd_dbg_is_rlc_restore_supported(pdd->dev))
                        amdgpu_gfx_off_ctrl(pdd->dev->adev, true);
 
+               /*
+                * Setting the debug flag in the trap handler requires that the TMA has been
+                * allocated, which occurs during CWSR initialization.
+                * In the event that CWSR has not been initialized at this point, setting the
+                * flag will be called again during CWSR initialization if the target process
+                * is still debug enabled.
+                */
+               kfd_process_set_trap_debug_flag(&pdd->qpd, true);
+
                if (!pdd->dev->kfd->shared_resources.enable_mes)
                        r = debug_refresh_runlist(pdd->dev->dqm);
                else
 
 void kfd_process_set_trap_handler(struct qcm_process_device *qpd,
                                  uint64_t tba_addr,
                                  uint64_t tma_addr);
+void kfd_process_set_trap_debug_flag(struct qcm_process_device *qpd,
+                                    bool enabled);
 
 /* CWSR initialization */
 int kfd_process_init_cwsr_apu(struct kfd_process *process, struct file *filep);
 
 
                memcpy(qpd->cwsr_kaddr, dev->kfd->cwsr_isa, dev->kfd->cwsr_isa_size);
 
+               kfd_process_set_trap_debug_flag(qpd, p->debug_trap_enabled);
+
                qpd->tma_addr = qpd->tba_addr + KFD_CWSR_TMA_OFFSET;
                pr_debug("set tba :0x%llx, tma:0x%llx, cwsr_kaddr:%p for pqm.\n",
                        qpd->tba_addr, qpd->tma_addr, qpd->cwsr_kaddr);
 
        memcpy(qpd->cwsr_kaddr, dev->kfd->cwsr_isa, dev->kfd->cwsr_isa_size);
 
+       kfd_process_set_trap_debug_flag(&pdd->qpd,
+                                       pdd->process->debug_trap_enabled);
+
        qpd->tma_addr = qpd->tba_addr + KFD_CWSR_TMA_OFFSET;
        pr_debug("set tba :0x%llx, tma:0x%llx, cwsr_kaddr:%p for pqm.\n",
                 qpd->tba_addr, qpd->tma_addr, qpd->cwsr_kaddr);
        return true;
 }
 
+void kfd_process_set_trap_debug_flag(struct qcm_process_device *qpd,
+                                    bool enabled)
+{
+       if (qpd->cwsr_kaddr) {
+               uint64_t *tma =
+                       (uint64_t *)(qpd->cwsr_kaddr + KFD_CWSR_TMA_OFFSET);
+               tma[2] = enabled;
+       }
+}
+
 /*
  * On return the kfd_process is fully operational and will be freed when the
  * mm is released