summaryrefslogtreecommitdiff
path: root/Source/Core/VideoCommon/Fifo.cpp
diff options
context:
space:
mode:
authorMarkus Wick <markus+github@selfnet.de>2015-04-06 13:24:17 +0200
committerMarkus Wick <markus+github@selfnet.de>2015-04-06 13:24:17 +0200
commit4669b50e2351830656e7d26ec2f0eec8b4740e69 (patch)
tree7da8e9583a95528ef48ade14003b8e381c600908 /Source/Core/VideoCommon/Fifo.cpp
parentea50dc240d0df56490aea86408dd0f621787950d (diff)
parent74795b45539ebee1c2af78fd6dca9022c3189c76 (diff)
Merge pull request #2172 from degasus/block_gpu_thread
Block gpu thread
Diffstat (limited to 'Source/Core/VideoCommon/Fifo.cpp')
-rw-r--r--Source/Core/VideoCommon/Fifo.cpp152
1 files changed, 94 insertions, 58 deletions
diff --git a/Source/Core/VideoCommon/Fifo.cpp b/Source/Core/VideoCommon/Fifo.cpp
index 8c7dbce49b..c5e49fb583 100644
--- a/Source/Core/VideoCommon/Fifo.cpp
+++ b/Source/Core/VideoCommon/Fifo.cpp
@@ -5,6 +5,7 @@
#include "Common/Atomic.h"
#include "Common/ChunkFile.h"
#include "Common/CPUDetect.h"
+#include "Common/Event.h"
#include "Common/FPURoundMode.h"
#include "Common/MemoryUtil.h"
#include "Common/Thread.h"
@@ -29,7 +30,6 @@ bool g_bSkipCurrentFrame = false;
static volatile bool GpuRunningState = false;
static volatile bool EmuRunningState = false;
-static std::mutex m_csHWVidOccupied;
// Most of this array is unlikely to be faulted in...
static u8 s_fifo_aux_data[FIFO_SIZE];
@@ -58,6 +58,12 @@ static u8* s_video_buffer_pp_read_ptr;
// polls, it's just atomic.
// - The pp_read_ptr is the CPU preprocessing version of the read_ptr.
+static Common::Flag s_gpu_is_running; // If this one is set, the gpu loop will be called at least once again
+static Common::Event s_gpu_new_work_event;
+
+static Common::Flag s_gpu_is_pending; // If this one is set, there might still be work to do
+static Common::Event s_gpu_done_event;
+
void Fifo_DoState(PointerWrap &p)
{
p.DoArray(s_video_buffer, FIFO_SIZE);
@@ -79,16 +85,12 @@ void Fifo_PauseAndLock(bool doLock, bool unpauseOnUnlock)
{
SyncGPU(SYNC_GPU_OTHER);
EmulatorState(false);
- if (!Core::IsGPUThread())
- m_csHWVidOccupied.lock();
- _dbg_assert_(COMMON, !CommandProcessor::fifo.isGpuReadingData);
+ FlushGpu();
}
else
{
if (unpauseOnUnlock)
EmulatorState(true);
- if (!Core::IsGPUThread())
- m_csHWVidOccupied.unlock();
}
}
@@ -127,17 +129,18 @@ void ExitGpuLoop()
{
// This should break the wait loop in CPU thread
CommandProcessor::fifo.bFF_GPReadEnable = false;
- SCPFifoStruct &fifo = CommandProcessor::fifo;
- while (fifo.isGpuReadingData)
- Common::YieldCPU();
+ FlushGpu();
+
// Terminate GPU thread loop
GpuRunningState = false;
EmuRunningState = true;
+ s_gpu_new_work_event.Set();
}
void EmulatorState(bool running)
{
EmuRunningState = running;
+ s_gpu_new_work_event.Set();
}
void SyncGPU(SyncGPUReason reason, bool may_move_read_ptr)
@@ -266,15 +269,10 @@ void ResetVideoBuffer()
// Purpose: Keep the Core HW updated about the CPU-GPU distance
void RunGpuLoop()
{
- std::lock_guard<std::mutex> lk(m_csHWVidOccupied);
GpuRunningState = true;
SCPFifoStruct &fifo = CommandProcessor::fifo;
u32 cyclesExecuted = 0;
- // If the host CPU has only two cores, idle loop instead of busy loop
- // This allows a system that we are maxing out in dual core mode to do other things
- bool yield_cpu = cpu_info.num_cores <= 2;
-
AsyncRequests::GetInstance()->SetEnable(true);
AsyncRequests::GetInstance()->SetPassthrough(false);
@@ -282,9 +280,10 @@ void RunGpuLoop()
{
g_video_backend->PeekMessages();
- AsyncRequests::GetInstance()->PullEvents();
- if (g_use_deterministic_gpu_thread)
+ if (g_use_deterministic_gpu_thread && EmuRunningState)
{
+ AsyncRequests::GetInstance()->PullEvents();
+
// All the fifo/CP stuff is on the CPU. We just need to run the opcode decoder.
u8* seen_ptr = s_video_buffer_seen_ptr;
u8* write_ptr = s_video_buffer_write_ptr;
@@ -300,17 +299,23 @@ void RunGpuLoop()
}
}
}
- else
+ else if (EmuRunningState)
{
+ AsyncRequests::GetInstance()->PullEvents();
+
CommandProcessor::SetCPStatusFromGPU();
- Common::AtomicStore(CommandProcessor::VITicks, CommandProcessor::m_cpClockOrigin);
+ if (!fifo.isGpuReadingData)
+ {
+ Common::AtomicStore(CommandProcessor::VITicks, CommandProcessor::m_cpClockOrigin);
+ }
+
+ bool run_loop = true;
// check if we are able to run this buffer
- while (GpuRunningState && EmuRunningState && !CommandProcessor::interruptWaiting && fifo.bFF_GPReadEnable && fifo.CPReadWriteDistance && !AtBreakpoint())
+ while (run_loop && !CommandProcessor::interruptWaiting && fifo.bFF_GPReadEnable && fifo.CPReadWriteDistance && !AtBreakpoint())
{
fifo.isGpuReadingData = true;
- CommandProcessor::isPossibleWaitingSetDrawDone = fifo.bFF_GPLinkEnable ? true : false;
if (!SConfig::GetInstance().m_LocalCoreStartupParameter.bSyncGPU || Common::AtomicLoad(CommandProcessor::VITicks) > CommandProcessor::m_cpClockOrigin)
{
@@ -338,6 +343,10 @@ void RunGpuLoop()
if ((write_ptr - s_video_buffer_read_ptr) == 0)
Common::AtomicStore(fifo.SafeCPReadPointer, fifo.CPReadPointer);
}
+ else
+ {
+ run_loop = false;
+ }
CommandProcessor::SetCPStatusFromGPU();
@@ -345,30 +354,28 @@ void RunGpuLoop()
// If we don't, s_swapRequested or s_efbAccessRequested won't be set to false
// leading the CPU thread to wait in Video_BeginField or Video_AccessEFB thus slowing things down.
AsyncRequests::GetInstance()->PullEvents();
- CommandProcessor::isPossibleWaitingSetDrawDone = false;
}
- fifo.isGpuReadingData = false;
+ // don't release the GPU running state on sync GPU waits
+ fifo.isGpuReadingData = !run_loop;
}
- if (EmuRunningState)
+ s_gpu_is_pending.Clear();
+ s_gpu_done_event.Set();
+
+ if (s_gpu_is_running.IsSet())
{
- // NOTE(jsd): Calling SwitchToThread() on Windows 7 x64 is a hot spot, according to profiler.
- // See https://docs.google.com/spreadsheet/ccc?key=0Ah4nh0yGtjrgdFpDeF9pS3V6RUotRVE3S3J4TGM1NlE#gid=0
- // for benchmark details.
- if (yield_cpu)
- Common::YieldCPU();
+ if (CommandProcessor::s_gpuMaySleep.IsSet())
+ {
+ // Reset the atomic flag. But as the CPU thread might have pushed some new data, we have to rerun the GPU loop
+ s_gpu_is_pending.Set();
+ s_gpu_is_running.Clear();
+ CommandProcessor::s_gpuMaySleep.Clear();
+ }
}
else
{
- // While the emu is paused, we still handle async requests then sleep.
- while (!EmuRunningState)
- {
- g_video_backend->PeekMessages();
- m_csHWVidOccupied.unlock();
- Common::SleepCurrentThread(1);
- m_csHWVidOccupied.lock();
- }
+ s_gpu_new_work_event.WaitFor(std::chrono::milliseconds(100));
}
}
// wake up SyncGPU if we were interrupted
@@ -377,6 +384,17 @@ void RunGpuLoop()
AsyncRequests::GetInstance()->SetPassthrough(true);
}
+void FlushGpu()
+{
+ if (!SConfig::GetInstance().m_LocalCoreStartupParameter.bCPUThread || g_use_deterministic_gpu_thread)
+ return;
+
+ while (s_gpu_is_running.IsSet() || s_gpu_is_pending.IsSet())
+ {
+ CommandProcessor::s_gpuMaySleep.Set();
+ s_gpu_done_event.Wait();
+ }
+}
bool AtBreakpoint()
{
@@ -386,41 +404,59 @@ bool AtBreakpoint()
void RunGpu()
{
- if (SConfig::GetInstance().m_LocalCoreStartupParameter.bCPUThread &&
- !g_use_deterministic_gpu_thread)
- return;
-
SCPFifoStruct &fifo = CommandProcessor::fifo;
- while (fifo.bFF_GPReadEnable && fifo.CPReadWriteDistance && !AtBreakpoint() )
+
+ // execute GPU
+ if (!SConfig::GetInstance().m_LocalCoreStartupParameter.bCPUThread || g_use_deterministic_gpu_thread)
{
- if (g_use_deterministic_gpu_thread)
+ bool reset_simd_state = false;
+ while (fifo.bFF_GPReadEnable && fifo.CPReadWriteDistance && !AtBreakpoint() )
{
- ReadDataFromFifoOnCPU(fifo.CPReadPointer);
+ if (g_use_deterministic_gpu_thread)
+ {
+ ReadDataFromFifoOnCPU(fifo.CPReadPointer);
+ }
+ else
+ {
+ if (!reset_simd_state)
+ {
+ FPURoundMode::SaveSIMDState();
+ FPURoundMode::LoadDefaultSIMDState();
+ reset_simd_state = true;
+ }
+ ReadDataFromFifo(fifo.CPReadPointer);
+ s_video_buffer_read_ptr = OpcodeDecoder_Run(DataReader(s_video_buffer_read_ptr, s_video_buffer_write_ptr), nullptr, false);
+ }
+
+ //DEBUG_LOG(COMMANDPROCESSOR, "Fifo wraps to base");
+
+ if (fifo.CPReadPointer == fifo.CPEnd)
+ fifo.CPReadPointer = fifo.CPBase;
+ else
+ fifo.CPReadPointer += 32;
+
+ fifo.CPReadWriteDistance -= 32;
}
- else
+ CommandProcessor::SetCPStatusFromGPU();
+
+ if (reset_simd_state)
{
- FPURoundMode::SaveSIMDState();
- FPURoundMode::LoadDefaultSIMDState();
- ReadDataFromFifo(fifo.CPReadPointer);
- s_video_buffer_read_ptr = OpcodeDecoder_Run(DataReader(s_video_buffer_read_ptr, s_video_buffer_write_ptr), nullptr, false);
FPURoundMode::LoadSIMDState();
}
+ }
- //DEBUG_LOG(COMMANDPROCESSOR, "Fifo wraps to base");
-
- if (fifo.CPReadPointer == fifo.CPEnd)
- fifo.CPReadPointer = fifo.CPBase;
- else
- fifo.CPReadPointer += 32;
-
- fifo.CPReadWriteDistance -= 32;
+ // wake up GPU thread
+ if (SConfig::GetInstance().m_LocalCoreStartupParameter.bCPUThread && !s_gpu_is_running.IsSet())
+ {
+ s_gpu_is_pending.Set();
+ s_gpu_is_running.Set();
+ s_gpu_new_work_event.Set();
}
- CommandProcessor::SetCPStatusFromGPU();
}
void Fifo_UpdateWantDeterminism(bool want)
{
- // We are paused (or not running at all yet) and have m_csHWVidOccupied, so
+ // We are paused (or not running at all yet), so
// it should be safe to change this.
const SCoreStartupParameter& param = SConfig::GetInstance().m_LocalCoreStartupParameter;
bool gpu_thread = false;