aboutsummaryrefslogtreecommitdiff
path: root/src/common/x64/cpu_wait.cpp
blob: 41d385f598706ea9f1d5b8a885cf03cc8bb78550 (plain) (blame)
1
2
3
4
5
6
7
8
9
10
11
12
13
14
15
16
17
18
19
20
21
22
23
24
25
26
27
28
29
30
31
32
33
34
35
36
37
38
39
40
41
42
43
44
45
46
47
48
49
50
51
52
53
54
55
56
57
58
59
60
61
62
63
64
65
66
67
68
69
70
71
72
73
74
75
// SPDX-FileCopyrightText: Copyright 2023 yuzu Emulator Project
// SPDX-License-Identifier: GPL-2.0-or-later

#include <thread>

#ifdef _MSC_VER
#include <intrin.h>
#endif

#include "common/x64/cpu_detect.h"
#include "common/x64/cpu_wait.h"
#include "common/x64/rdtsc.h"

namespace Common::X64 {

namespace {

// 100,000 cycles is a reasonable amount of time to wait to save on CPU resources.
// For reference:
// At 1 GHz, 100K cycles is 100us
// At 2 GHz, 100K cycles is 50us
// At 4 GHz, 100K cycles is 25us
constexpr auto PauseCycles = 100'000U;

} // Anonymous namespace

#ifdef _MSC_VER
__forceinline static void TPAUSE() {
    static constexpr auto RequestC02State = 0U;
    _tpause(RequestC02State, FencedRDTSC() + PauseCycles);
}

__forceinline static void MWAITX() {
    static constexpr auto EnableWaitTimeFlag = 1U << 1;
    static constexpr auto RequestC1State = 0U;

    // monitor_var should be aligned to a cache line.
    alignas(64) u64 monitor_var{};
    _mm_monitorx(&monitor_var, 0, 0);
    _mm_mwaitx(EnableWaitTimeFlag, RequestC1State, PauseCycles);
}
#else
static void TPAUSE() {
    static constexpr auto RequestC02State = 0U;
    const auto tsc = FencedRDTSC() + PauseCycles;
    const auto eax = static_cast<u32>(tsc & 0xFFFFFFFF);
    const auto edx = static_cast<u32>(tsc >> 32);
    asm volatile("tpause %0" : : "r"(RequestC02State), "d"(edx), "a"(eax));
}

static void MWAITX() {
    static constexpr auto EnableWaitTimeFlag = 1U << 1;
    static constexpr auto RequestC1State = 0U;

    // monitor_var should be aligned to a cache line.
    alignas(64) u64 monitor_var{};
    asm volatile("monitorx" : : "a"(&monitor_var), "c"(0), "d"(0));
    asm volatile("mwaitx" : : "a"(RequestC1State), "b"(PauseCycles), "c"(EnableWaitTimeFlag));
}
#endif

void MicroSleep() {
    static const bool has_waitpkg = GetCPUCaps().waitpkg;
    static const bool has_monitorx = GetCPUCaps().monitorx;

    if (has_waitpkg) {
        TPAUSE();
    } else if (has_monitorx) {
        MWAITX();
    } else {
        std::this_thread::yield();
    }
}

} // namespace Common::X64