all: add support for multicore scheduler

This commit adds support for a scheduler that runs a scheduler on all
available cores. It is meant to be used on baremetal systems with a
fixed number of cores, such as the RP2040.

The initial implementation adds support for multicore scheduling to the
riscv-qemu target as a convenient testing target. This means that this
new multicore scheduler is tested in CI, including a bunch of standard
library tests (`make tinygo-test-baremetal`). This should ensure the new
scheduler is reasonably well tested before trying to use it on
harder-to-debug targets like the RP2040.
This commit is contained in:
Ayke van Laethem
2025-04-03 13:29:55 +02:00
committed by Ron Evans
parent 0c7c2926f9
commit 60f8a62978
35 changed files with 1237 additions and 160 deletions
+1 -1
View File
@@ -490,7 +490,7 @@ func loadProgramSize(path string, packagePathMap map[string]string) (*programSiz
continue
}
if section.Type == elf.SHT_NOBITS {
if section.Name == ".stack" {
if strings.HasPrefix(section.Name, ".stack") {
// TinyGo emits stack sections on microcontroller using the
// ".stack" name.
// This is a bit ugly, but I don't think there is a way to
+1 -1
View File
@@ -42,7 +42,7 @@ func TestBinarySize(t *testing.T) {
// This is a small number of very diverse targets that we want to test.
tests := []sizeTest{
// microcontrollers
{"hifive1b", "examples/echo", 4556, 280, 0, 2268},
{"hifive1b", "examples/echo", 4556, 280, 0, 2264},
{"microbit", "examples/serial", 2920, 388, 8, 2272},
{"wioterminal", "examples/pininterrupt", 7379, 1489, 116, 6912},
+5
View File
@@ -110,6 +110,11 @@ func (c *Config) BuildTags() []string {
"math_big_pure_go", // to get math/big to work
"gc." + c.GC(), "scheduler." + c.Scheduler(), // used inside the runtime package
"serial." + c.Serial()}...) // used inside the machine package
switch c.Scheduler() {
case "threads", "cores":
default:
tags = append(tags, "tinygo.unicore")
}
for i := 1; i <= c.GoMinorVersion; i++ {
tags = append(tags, fmt.Sprintf("go1.%d", i))
}
+1 -1
View File
@@ -10,7 +10,7 @@ import (
var (
validBuildModeOptions = []string{"default", "c-shared", "wasi-legacy"}
validGCOptions = []string{"none", "leaking", "conservative", "custom", "precise", "boehm"}
validSchedulerOptions = []string{"none", "tasks", "asyncify", "threads"}
validSchedulerOptions = []string{"none", "tasks", "asyncify", "threads", "cores"}
validSerialOptions = []string{"none", "uart", "usb", "rtt"}
validPrintSizeOptions = []string{"none", "short", "full", "html"}
validPanicStrategyOptions = []string{"print", "trap"}
+1 -1
View File
@@ -10,7 +10,7 @@ import (
func TestVerifyOptions(t *testing.T) {
expectedGCError := errors.New(`invalid gc option 'incorrect': valid values are none, leaking, conservative, custom, precise, boehm`)
expectedSchedulerError := errors.New(`invalid scheduler option 'incorrect': valid values are none, tasks, asyncify, threads`)
expectedSchedulerError := errors.New(`invalid scheduler option 'incorrect': valid values are none, tasks, asyncify, threads, cores`)
expectedPrintSizeError := errors.New(`invalid size option 'incorrect': valid values are none, short, full, html`)
expectedPanicStrategyError := errors.New(`invalid panic option 'incorrect': valid values are print, trap`)
+39
View File
@@ -3,8 +3,47 @@
.type _start,@function
_start:
// If we're on a multicore system, we need to wait for hart 0 to wake us up.
#if TINYGO_CORES > 1
csrr a0, mhartid
// Hart 0 stack
bnez a0, 1f
la sp, _stack_top
1:
// Hart 1 stack
li a1, 1
bne a0, a1, 2f
la sp, _stack1_top
2:
// Hart 2 stack
#if TINYGO_CORES >= 3
li a1, 2
bne a0, a1, 3f
la sp, _stack2_top
#endif
3:
// Hart 3 stack
#if TINYGO_CORES >= 4
li a1, 3
bne a0, a1, 4f
la sp, _stack3_top
#endif
4:
// done
#if TINYGO_CORES > 4
#error only up to 4 cores are supported at the moment!
#endif
#else
// Load the stack pointer.
la sp, _stack_top
#endif
// Load the globals pointer. The program will load pointers relative to this
// register, so it must be set to the right value on startup.
+1 -1
View File
@@ -1,4 +1,4 @@
//go:build !scheduler.threads
//go:build tinygo.unicore
package task
+1 -1
View File
@@ -1,4 +1,4 @@
//go:build scheduler.threads
//go:build !tinygo.unicore
package task
+1 -1
View File
@@ -1,4 +1,4 @@
//go:build !scheduler.threads
//go:build tinygo.unicore
package task
+64
View File
@@ -0,0 +1,64 @@
//go:build scheduler.cores
package task
import "runtime/interrupt"
// A futex is a way for userspace to wait with the pointer as the key, and for
// another thread to wake one or all waiting threads keyed on the same pointer.
//
// A futex does not change the underlying value, it only reads it before to prevent
// lost wake-ups.
type Futex struct {
Uint32
waiters Stack
}
// Atomically check for cmp to still be equal to the futex value and if so, go
// to sleep. Return true if we were definitely awoken by a call to Wake or
// WakeAll, and false if we can't be sure of that.
func (f *Futex) Wait(cmp uint32) (awoken bool) {
mask := lockFutex()
if f.Uint32.Load() != cmp {
unlockFutex(mask)
return false
}
// Push the current goroutine onto the waiter stack.
f.waiters.Push(Current())
unlockFutex(mask)
// Pause until this task is awoken by Wake/WakeAll.
Pause()
// We were awoken by a call to Wake or WakeAll. There is no chance for
// spurious wakeups.
return true
}
// Wake a single waiter.
func (f *Futex) Wake() {
mask := lockFutex()
if t := f.waiters.Pop(); t != nil {
scheduleTask(t)
}
unlockFutex(mask)
}
// Wake all waiters.
func (f *Futex) WakeAll() {
mask := lockFutex()
for t := f.waiters.Pop(); t != nil; t = f.waiters.Pop() {
scheduleTask(t)
}
unlockFutex(mask)
}
//go:linkname lockFutex runtime.lockFutex
func lockFutex() interrupt.State
//go:linkname unlockFutex runtime.unlockFutex
func unlockFutex(interrupt.State)
+1 -1
View File
@@ -1,4 +1,4 @@
//go:build !scheduler.threads
//go:build tinygo.unicore
package task
+1 -1
View File
@@ -1,4 +1,4 @@
//go:build scheduler.threads
//go:build !tinygo.unicore
package task
+1 -1
View File
@@ -1,4 +1,4 @@
//go:build !scheduler.threads
//go:build tinygo.unicore
package task
+1 -1
View File
@@ -1,4 +1,4 @@
//go:build scheduler.threads
//go:build !tinygo.unicore
package task
+30 -17
View File
@@ -12,9 +12,9 @@ type Queue struct {
// Push a task onto the queue.
func (q *Queue) Push(t *Task) {
i := interrupt.Disable()
mask := lockAtomics()
if asserts && t.Next != nil {
interrupt.Restore(i)
unlockAtomics(mask)
panic("runtime: pushing a task to a queue with a non-nil Next pointer")
}
if q.tail != nil {
@@ -25,15 +25,15 @@ func (q *Queue) Push(t *Task) {
if q.head == nil {
q.head = t
}
interrupt.Restore(i)
unlockAtomics(mask)
}
// Pop a task off of the queue.
func (q *Queue) Pop() *Task {
i := interrupt.Disable()
mask := lockAtomics()
t := q.head
if t == nil {
interrupt.Restore(i)
unlockAtomics(mask)
return nil
}
q.head = t.Next
@@ -41,13 +41,13 @@ func (q *Queue) Pop() *Task {
q.tail = nil
}
t.Next = nil
interrupt.Restore(i)
unlockAtomics(mask)
return t
}
// Append pops the contents of another queue and pushes them onto the end of this queue.
func (q *Queue) Append(other *Queue) {
i := interrupt.Disable()
mask := lockAtomics()
if q.head == nil {
q.head = other.head
} else {
@@ -55,14 +55,14 @@ func (q *Queue) Append(other *Queue) {
}
q.tail = other.tail
other.head, other.tail = nil, nil
interrupt.Restore(i)
unlockAtomics(mask)
}
// Empty checks if the queue is empty.
func (q *Queue) Empty() bool {
i := interrupt.Disable()
mask := lockAtomics()
empty := q.head == nil
interrupt.Restore(i)
unlockAtomics(mask)
return empty
}
@@ -75,24 +75,24 @@ type Stack struct {
// Push a task onto the stack.
func (s *Stack) Push(t *Task) {
i := interrupt.Disable()
mask := lockAtomics()
if asserts && t.Next != nil {
interrupt.Restore(i)
unlockAtomics(mask)
panic("runtime: pushing a task to a stack with a non-nil Next pointer")
}
s.top, t.Next = t, s.top
interrupt.Restore(i)
unlockAtomics(mask)
}
// Pop a task off of the stack.
func (s *Stack) Pop() *Task {
i := interrupt.Disable()
mask := lockAtomics()
t := s.top
if t != nil {
s.top = t.Next
t.Next = nil
}
interrupt.Restore(i)
unlockAtomics(mask)
return t
}
@@ -112,13 +112,26 @@ func (t *Task) tail() *Task {
// Queue moves the contents of the stack into a queue.
// Elements can be popped from the queue in the same order that they would be popped from the stack.
func (s *Stack) Queue() Queue {
i := interrupt.Disable()
mask := lockAtomics()
head := s.top
s.top = nil
q := Queue{
head: head,
tail: head.tail(),
}
interrupt.Restore(i)
unlockAtomics(mask)
return q
}
// Use runtime.lockAtomics and runtime.unlockAtomics so that Queue and Stack
// work correctly even on multicore systems. These functions are normally used
// to implement atomic operations, but the same spinlock can also be used for
// Queue/Stack operations which are very fast.
// These functions are just plain old interrupt disable/restore on non-multicore
// systems.
//go:linkname lockAtomics runtime.lockAtomics
func lockAtomics() interrupt.State
//go:linkname unlockAtomics runtime.unlockAtomics
func unlockAtomics(mask interrupt.State)
+17
View File
@@ -24,11 +24,28 @@ type Task struct {
// This is needed for some crypto packages.
FipsIndicator uint8
// State of the goroutine: running, paused, or must-resume-next-pause.
// This extra field doesn't increase memory usage on 32-bit CPUs and above,
// since it falls into the padding of the FipsIndicator bit above.
RunState uint8
// DeferFrame stores a pointer to the (stack allocated) defer frame of the
// goroutine that is used for the recover builtin.
DeferFrame unsafe.Pointer
}
const (
// Initial state: the goroutine state is saved on the stack.
RunStatePaused = iota
// The goroutine is running right now.
RunStateRunning
// The goroutine is running, but already marked as "can resume".
// The next call to Pause() won't actually pause the goroutine.
RunStateResuming
)
// DataUint32 returns the Data field as a uint32. The value is only valid after
// setting it through SetDataUint32 or by storing to it using DataAtomicUint32.
func (t *Task) DataUint32() uint32 {
+1 -34
View File
@@ -1,9 +1,8 @@
//go:build scheduler.tasks
//go:build scheduler.tasks || scheduler.cores
package task
import (
"runtime/interrupt"
"unsafe"
)
@@ -32,44 +31,12 @@ type state struct {
canaryPtr *uintptr
}
// currentTask is the current running task, or nil if currently in the scheduler.
var currentTask *Task
// Current returns the current active task.
func Current() *Task {
return currentTask
}
// Pause suspends the current task and returns to the scheduler.
// This function may only be called when running on a goroutine stack, not when running on the system stack or in an interrupt.
func Pause() {
// Check whether the canary (the lowest address of the stack) is still
// valid. If it is not, a stack overflow has occurred.
if *currentTask.state.canaryPtr != stackCanary {
runtimePanic("goroutine stack overflow")
}
if interrupt.In() {
runtimePanic("blocked inside interrupt")
}
currentTask.state.pause()
}
//export tinygo_task_exit
func taskExit() {
// TODO: explicitly free the stack after switching back to the scheduler.
Pause()
}
// Resume the task until it pauses or completes.
// This may only be called from the scheduler.
func (t *Task) Resume() {
currentTask = t
t.gcData.swap()
t.state.resume()
t.gcData.swap()
currentTask = nil
}
// initialize the state and prepare to call the specified function with the specified argument bundle.
func (s *state) initialize(fn uintptr, args unsafe.Pointer, stackSize uintptr) {
// Create a stack.
+53
View File
@@ -0,0 +1,53 @@
//go:build scheduler.cores
package task
import "runtime/interrupt"
// Current returns the current active task.
//
//go:linkname Current runtime.currentTask
func Current() *Task
// Pause suspends the current task and returns to the scheduler.
// This function may only be called when running on a goroutine stack, not when running on the system stack or in an interrupt.
func Pause() {
lockScheduler()
PauseLocked()
}
// PauseLocked is the same as Pause, but must be called with the scheduler lock
// already taken.
func PauseLocked() {
// Check whether the canary (the lowest address of the stack) is still
// valid. If it is not, a stack overflow has occurred.
current := Current()
if *current.state.canaryPtr != stackCanary {
runtimePanic("goroutine stack overflow")
}
if interrupt.In() {
runtimePanic("blocked inside interrupt")
}
if current.RunState == RunStateResuming {
// Another core already marked this goroutine as ready to resume.
current.RunState = RunStateRunning
unlockScheduler()
return
}
current.RunState = RunStatePaused
current.state.pause()
}
// Resume the task until it pauses or completes.
// This may only be called from the scheduler.
func (t *Task) Resume() {
t.gcData.swap()
t.state.resume()
t.gcData.swap()
}
//go:linkname lockScheduler runtime.lockScheduler
func lockScheduler()
//go:linkname unlockScheduler runtime.unlockScheduler
func unlockScheduler()
+13 -6
View File
@@ -1,10 +1,16 @@
//go:build scheduler.tasks && tinygo.riscv
//go:build (scheduler.tasks || scheduler.cores) && tinygo.riscv
package task
import "unsafe"
var systemStack uintptr
// Returns a pointer where the system stack can be stored.
// This is a layering violation! We should probably refactor this so that we
// don't need such gymnastics to store the system stack pointer. (It should
// probably be moved to the runtime).
//
//go:linkname runtime_systemStackPtr runtime.systemStackPtr
func runtime_systemStackPtr() *uintptr
// calleeSavedRegs is the list of registers that must be saved and restored when
// switching between tasks. Also see scheduler_riscv.S that relies on the
@@ -50,17 +56,18 @@ func (s *state) archInit(r *calleeSavedRegs, fn uintptr, args unsafe.Pointer) {
}
func (s *state) resume() {
swapTask(s.sp, &systemStack)
swapTask(s.sp, runtime_systemStackPtr())
}
func (s *state) pause() {
newStack := systemStack
systemStack = 0
systemStackPtr := runtime_systemStackPtr()
newStack := *systemStackPtr
*systemStackPtr = 0
swapTask(newStack, &s.sp)
}
// SystemStack returns the system stack pointer when called from a task stack.
// When called from the system stack, it returns 0.
func SystemStack() uintptr {
return systemStack
return *runtime_systemStackPtr()
}
+37
View File
@@ -0,0 +1,37 @@
//go:build scheduler.tasks
package task
import "runtime/interrupt"
// currentTask is the current running task, or nil if currently in the scheduler.
var currentTask *Task
// Current returns the current active task.
func Current() *Task {
return currentTask
}
// Pause suspends the current task and returns to the scheduler.
// This function may only be called when running on a goroutine stack, not when running on the system stack or in an interrupt.
func Pause() {
// Check whether the canary (the lowest address of the stack) is still
// valid. If it is not, a stack overflow has occurred.
if *currentTask.state.canaryPtr != stackCanary {
runtimePanic("goroutine stack overflow")
}
if interrupt.In() {
runtimePanic("blocked inside interrupt")
}
currentTask.state.pause()
}
// Resume the task until it pauses or completes.
// This may only be called from the scheduler.
func (t *Task) Resume() {
currentTask = t
t.gcData.swap()
t.state.resume()
t.gcData.swap()
currentTask = nil
}
+30 -31
View File
@@ -6,7 +6,6 @@
package runtime
import (
"runtime/interrupt"
_ "unsafe"
)
@@ -23,27 +22,27 @@ import (
func __atomic_load_2(ptr *uint16, ordering uintptr) uint16 {
// The LLVM docs for this say that there is a val argument after the pointer.
// That is a typo, and the GCC docs omit it.
mask := interrupt.Disable()
mask := lockAtomics()
val := *ptr
interrupt.Restore(mask)
unlockAtomics(mask)
return val
}
//export __atomic_store_2
func __atomic_store_2(ptr *uint16, val uint16, ordering uintptr) {
mask := interrupt.Disable()
mask := lockAtomics()
*ptr = val
interrupt.Restore(mask)
unlockAtomics(mask)
}
//go:inline
func doAtomicCAS16(ptr *uint16, expected, desired uint16) uint16 {
mask := interrupt.Disable()
mask := lockAtomics()
old := *ptr
if old == expected {
*ptr = desired
}
interrupt.Restore(mask)
unlockAtomics(mask)
return old
}
@@ -61,10 +60,10 @@ func __atomic_compare_exchange_2(ptr, expected *uint16, desired uint16, successO
//go:inline
func doAtomicSwap16(ptr *uint16, new uint16) uint16 {
mask := interrupt.Disable()
mask := lockAtomics()
old := *ptr
*ptr = new
interrupt.Restore(mask)
unlockAtomics(mask)
return old
}
@@ -80,11 +79,11 @@ func __atomic_exchange_2(ptr *uint16, new uint16, ordering uintptr) uint16 {
//go:inline
func doAtomicAdd16(ptr *uint16, value uint16) (old, new uint16) {
mask := interrupt.Disable()
mask := lockAtomics()
old = *ptr
new = old + value
*ptr = new
interrupt.Restore(mask)
unlockAtomics(mask)
return old, new
}
@@ -112,27 +111,27 @@ func __atomic_add_fetch_2(ptr *uint16, value uint16, ordering uintptr) uint16 {
func __atomic_load_4(ptr *uint32, ordering uintptr) uint32 {
// The LLVM docs for this say that there is a val argument after the pointer.
// That is a typo, and the GCC docs omit it.
mask := interrupt.Disable()
mask := lockAtomics()
val := *ptr
interrupt.Restore(mask)
unlockAtomics(mask)
return val
}
//export __atomic_store_4
func __atomic_store_4(ptr *uint32, val uint32, ordering uintptr) {
mask := interrupt.Disable()
mask := lockAtomics()
*ptr = val
interrupt.Restore(mask)
unlockAtomics(mask)
}
//go:inline
func doAtomicCAS32(ptr *uint32, expected, desired uint32) uint32 {
mask := interrupt.Disable()
mask := lockAtomics()
old := *ptr
if old == expected {
*ptr = desired
}
interrupt.Restore(mask)
unlockAtomics(mask)
return old
}
@@ -150,10 +149,10 @@ func __atomic_compare_exchange_4(ptr, expected *uint32, desired uint32, successO
//go:inline
func doAtomicSwap32(ptr *uint32, new uint32) uint32 {
mask := interrupt.Disable()
mask := lockAtomics()
old := *ptr
*ptr = new
interrupt.Restore(mask)
unlockAtomics(mask)
return old
}
@@ -169,11 +168,11 @@ func __atomic_exchange_4(ptr *uint32, new uint32, ordering uintptr) uint32 {
//go:inline
func doAtomicAdd32(ptr *uint32, value uint32) (old, new uint32) {
mask := interrupt.Disable()
mask := lockAtomics()
old = *ptr
new = old + value
*ptr = new
interrupt.Restore(mask)
unlockAtomics(mask)
return old, new
}
@@ -201,27 +200,27 @@ func __atomic_add_fetch_4(ptr *uint32, value uint32, ordering uintptr) uint32 {
func __atomic_load_8(ptr *uint64, ordering uintptr) uint64 {
// The LLVM docs for this say that there is a val argument after the pointer.
// That is a typo, and the GCC docs omit it.
mask := interrupt.Disable()
mask := lockAtomics()
val := *ptr
interrupt.Restore(mask)
unlockAtomics(mask)
return val
}
//export __atomic_store_8
func __atomic_store_8(ptr *uint64, val uint64, ordering uintptr) {
mask := interrupt.Disable()
mask := lockAtomics()
*ptr = val
interrupt.Restore(mask)
unlockAtomics(mask)
}
//go:inline
func doAtomicCAS64(ptr *uint64, expected, desired uint64) uint64 {
mask := interrupt.Disable()
mask := lockAtomics()
old := *ptr
if old == expected {
*ptr = desired
}
interrupt.Restore(mask)
unlockAtomics(mask)
return old
}
@@ -239,10 +238,10 @@ func __atomic_compare_exchange_8(ptr, expected *uint64, desired uint64, successO
//go:inline
func doAtomicSwap64(ptr *uint64, new uint64) uint64 {
mask := interrupt.Disable()
mask := lockAtomics()
old := *ptr
*ptr = new
interrupt.Restore(mask)
unlockAtomics(mask)
return old
}
@@ -258,11 +257,11 @@ func __atomic_exchange_8(ptr *uint64, new uint64, ordering uintptr) uint64 {
//go:inline
func doAtomicAdd64(ptr *uint64, value uint64) (old, new uint64) {
mask := interrupt.Disable()
mask := lockAtomics()
old = *ptr
new = old + value
*ptr = new
interrupt.Restore(mask)
unlockAtomics(mask)
return old, new
}
+101
View File
@@ -0,0 +1,101 @@
//go:build scheduler.cores
package runtime
import (
"internal/task"
"sync/atomic"
)
// Normally 0. During a GC scan it has various purposes for signalling between
// the core running the GC and the other cores in the system.
var gcScanState atomic.Uint32
// Start GC scan by pausing the world (all other cores) and scanning their
// stacks. It doesn't resume the world.
func gcMarkReachable() {
core := currentCPU()
// Interrupt all other cores.
gcScanState.Store(1)
for i := uint32(0); i < numCPU; i++ {
if i == core {
continue
}
gcPauseCore(i)
}
// Scan the stack(s) of the current core.
scanCurrentStack()
if !task.OnSystemStack() {
// Mark system stack.
markRoots(task.SystemStack(), coreStackTop(core))
}
// Scan globals.
findGlobals(markRoots)
// Busy-wait until all the other cores are ready. They certainly should be,
// after the scanning we did above.
for gcScanState.Load() != numCPU {
spinLoopHint()
}
gcScanState.Store(0)
// Signal each core in turn that they can scan the stack.
for i := uint32(0); i < numCPU; i++ {
if i == core {
continue
}
// Wake up the core to scan the stack.
gcSignalCore(i)
// Busy-wait until this core finished scanning.
for gcScanState.Load() == 0 {
spinLoopHint()
}
gcScanState.Store(0)
}
// All the stack are now scanned.
}
//go:export tinygo_scanCurrentStack
func scanCurrentStack()
//go:export tinygo_scanstack
func scanstack(sp uintptr) {
// Mark the current stack.
// This function is called by scanCurrentStack, after pushing all registers
// onto the stack.
if task.OnSystemStack() {
// This is the system stack.
// Scan all words on the stack.
markRoots(sp, coreStackTop(currentCPU()))
} else {
// This is a goroutine stack.
markCurrentGoroutineStack(sp)
}
}
// Resume the world after a call to gcMarkReachable.
func gcResumeWorld() {
// Signal each core that they can resume.
hartID := currentCPU()
for i := uint32(0); i < numCPU; i++ {
if i == hartID {
continue
}
// Signal the core.
gcSignalCore(i)
}
// Busy-wait until the core acknowledges the signal (and is going to return
// from the interrupt handler).
for gcScanState.Load() != numCPU-1 {
spinLoopHint()
}
gcScanState.Store(0)
}
+8 -2
View File
@@ -1,8 +1,14 @@
//go:build (gc.conservative || gc.precise || gc.boehm) && !tinygo.wasm && !scheduler.threads
//go:build (gc.conservative || gc.precise || gc.boehm) && !tinygo.wasm && !scheduler.threads && !scheduler.cores
package runtime
import "internal/task"
import (
"internal/task"
"sync/atomic"
)
// Unused.
var gcScanState atomic.Uint32
func gcMarkReachable() {
markStack()
+2 -1
View File
@@ -99,7 +99,8 @@ func runtimePanicAt(addr unsafe.Pointer, msg string) {
} else {
printstring("panic: runtime error: ")
}
println(msg)
printstring(msg)
printnl()
abort()
}
-13
View File
@@ -1,7 +1,6 @@
package runtime
import (
"internal/task"
"unsafe"
)
@@ -9,18 +8,6 @@ type stringer interface {
String() string
}
// Lock to make sure print calls do not interleave.
// This is a no-op lock on systems that do not have parallelism.
var printLock task.PMutex
func printlock() {
printLock.Lock()
}
func printunlock() {
printLock.Unlock()
}
//go:nobounds
func printstring(s string) {
for i := 0; i < len(s); i++ {
+380 -27
View File
@@ -4,26 +4,82 @@ package runtime
import (
"device/riscv"
"internal/task"
"math/bits"
"runtime/interrupt"
"runtime/volatile"
"sync/atomic"
"unsafe"
)
// This file implements the VirtIO RISC-V interface implemented in QEMU, which
// is an interface designed for emulation.
const numCPU = 4
//export main
func main() {
preinit()
// Set the interrupt address.
// Note that this address must be aligned specially, otherwise the MODE bits
// of MTVEC won't be zero.
riscv.MTVEC.Set(uintptr(unsafe.Pointer(&handleInterruptASM)))
// Enable software interrupts. We'll need them to wake up other cores.
riscv.MIE.SetBits(riscv.MIE_MSIE)
// If we're not hart 0, wait until we get the signal everything has been set
// up.
if hartID := riscv.MHARTID.Get(); hartID != 0 {
// Wait until we get the signal this hart is ready to start.
// Note that interrupts are disabled, which means that the interrupt
// isn't actually taken. But we can still wait for it using wfi.
// If the cores scheduler is not used, we'll stay in this state forever.
for riscv.MIP.Get()&riscv.MIP_MSIP == 0 {
riscv.Asm("wfi")
}
// Clear the software interrupt.
aclintMSWI.MSIP[hartID].Set(0)
// Now that we've cleared the software interrupt, we can enable
// interrupts as was already done on hart 0.
riscv.MSTATUS.SetBits(riscv.MSTATUS_MIE)
// Also enable timer interrupts, for sleepTicksMulticore.
riscv.MIE.SetBits(riscv.MIE_MTIE)
// Now start running the scheduler on this core.
schedulerLock.Lock()
scheduler(false)
// The scheduler exited, which means main returned and the program
// should exit immediately.
// Signal hart 0 to exit.
exitCodePlusOne.Store(0 + 1) // exit code 0
aclintMSWI.MSIP[0].Set(1)
// Unlock the scheduler to be sure. Shouldn't be needed.
schedulerLock.Unlock()
// Wait until hart 0 actually exits.
for {
riscv.Asm("wfi")
}
}
// Enable global interrupts now that they've been set up.
// This is currently only for timer interrupts.
riscv.MSTATUS.SetBits(riscv.MSTATUS_MIE)
// Set all MTIMECMP registers to a value that clears the MTIP bit in MIP.
// If we don't do this, the wfi instruction won't work as expected.
for i := 0; i < numCPU; i++ {
aclintMTIMECMP[i].Set(0xffff_ffff_ffff_ffff)
}
// Enable timer interrupts on hart 0.
riscv.MIE.SetBits(riscv.MIE_MTIE)
run()
exit(0)
}
@@ -37,13 +93,32 @@ func handleInterrupt() {
code := uint(cause &^ (1 << 31))
if cause&(1<<31) != 0 {
// Topmost bit is set, which means that it is an interrupt.
hartID := currentCPU()
switch code {
case riscv.MachineSoftwareInterrupt:
if exitCodePlusOne.Load() != 0 {
exitNow(exitCodePlusOne.Load() - 1)
}
if gcScanState.Load() != 0 {
// The GC needs to run.
gcInterruptHandler(hartID)
}
checkpoint := &schedulerWaitCheckpoints[hartID]
if checkpoint.Saved() {
aclintMSWI.MSIP[hartID].Set(0)
riscv.MCAUSE.Set(0)
checkpoint.Jump()
}
case riscv.MachineTimerInterrupt:
// Signal timeout.
timerWakeup.Set(1)
// Disable the timer, to avoid triggering the interrupt right after
// this interrupt returns.
riscv.MIE.ClearBits(riscv.MIE_MTIE)
if sleepCheckpoint.Saved() {
// Set MTIMECMP to a high value so that MTIP goes low.
aclintMTIMECMP[hartID].Set(0xffff_ffff_ffff_ffff)
riscv.MCAUSE.Set(0)
sleepCheckpoint.Jump()
}
default:
runtimePanic("unknown interrupt")
abort()
}
} else {
// Topmost bit is clear, so it is an exception of some sort.
@@ -57,6 +132,79 @@ func handleInterrupt() {
riscv.MCAUSE.Set(0)
}
// The GC interrupted this core for the stop-the-world phase.
// This function handles that, and only returns after the stop-the-world phase
// ended.
func gcInterruptHandler(hartID uint32) {
// *only* enable the MSIE interrupt
savedMIE := riscv.MIE.Get()
riscv.MIE.Set(riscv.MIE_MSIE)
// Disable this interrupt (to be enabled again soon).
aclintMSWI.MSIP[hartID].Set(0)
// Let the GC know we're ready.
gcScanState.Add(1)
// Wait until we get a signal to start scanning.
for riscv.MIP.Get()&riscv.MIP_MSIP == 0 {
riscv.Asm("wfi")
}
aclintMSWI.MSIP[hartID].Set(0)
// Scan the stack(s) of this core.
scanCurrentStack()
if !task.OnSystemStack() {
// Mark system stack.
markRoots(task.SystemStack(), coreStackTop(hartID))
}
// Signal we've finished scanning.
gcScanState.Store(1)
// Wait until we get a signal that the stop-the-world phase has ended.
for riscv.MIP.Get()&riscv.MIP_MSIP == 0 {
riscv.Asm("wfi")
}
aclintMSWI.MSIP[hartID].Set(0)
// Restore MIE bits.
riscv.MIE.Set(savedMIE)
// Signal we received the signal and are going to exit the interrupt.
gcScanState.Add(1)
}
//go:extern _stack_top
var stack0TopSymbol [0]byte
//go:extern _stack1_top
var stack1TopSymbol [0]byte
//go:extern _stack2_top
var stack2TopSymbol [0]byte
//go:extern _stack3_top
var stack3TopSymbol [0]byte
// Returns the stack top (highest address) of the system stack of the given
// core.
func coreStackTop(core uint32) uintptr {
switch core {
case 0:
return uintptr(unsafe.Pointer(&stack0TopSymbol))
case 1:
return uintptr(unsafe.Pointer(&stack1TopSymbol))
case 2:
return uintptr(unsafe.Pointer(&stack2TopSymbol))
case 3:
return uintptr(unsafe.Pointer(&stack3TopSymbol))
default:
runtimePanic("unexpected core")
return 0
}
}
// One tick is 100ns by default in QEMU.
// (This is not a standard, just the default used by QEMU).
func ticksToNanoseconds(ticks timeUnit) int64 {
@@ -67,22 +215,79 @@ func nanosecondsToTicks(ns int64) timeUnit {
return timeUnit(ns / 100) // one tick is 100ns
}
var timerWakeup volatile.Register8
var sleepCheckpoint interrupt.Checkpoint
func sleepTicks(d timeUnit) {
// Enable the timer.
target := uint64(ticks() + d)
aclintMTIMECMP.Set(target)
riscv.MIE.SetBits(riscv.MIE_MTIE)
hartID := currentCPU()
if sleepCheckpoint.Save() {
// Configure timeout.
target := uint64(ticks() + d)
aclintMTIMECMP[hartID].Set(target)
// Wait until it fires.
for {
if timerWakeup.Get() != 0 {
timerWakeup.Set(0)
// Disable timer.
break
// Wait for the interrupt to happen.
for {
riscv.Asm("wfi")
}
}
// We got awoken.
}
// Currently sleeping core, or 0xff.
// Must only be accessed with the scheduler lock held.
var sleepingCore uint8 = 0xff
// Return whether another core is sleeping.
// May only be called with the scheduler lock held.
func hasSleepingCore() bool {
return sleepingCore != 0xff
}
// Almost identical to sleepTicks, except that it will unlock/lock the scheduler
// while sleeping and is interruptible by interruptSleepTicksMulticore.
// This may only be called with the scheduler lock held.
func sleepTicksMulticore(d timeUnit) {
// Disable interrupts while configuring sleep.
// This is needed because unlocking the scheduler and setting the timer
// interrupt need to happen atomically.
riscv.MSTATUS.ClearBits(riscv.MSTATUS_MIE)
hartID := currentCPU()
if sleepCheckpoint.Save() {
sleepingCore = uint8(hartID)
// Configure timeout.
target := uint64(ticks() + d)
aclintMTIMECMP[hartID].Set(target)
// Unlock, now that the timeout has been set (so that
// interruptSleepTicksMulticore will see the correct wakeup time).
schedulerLock.Unlock()
// Sleep has been configured, interrupts may happen again.
riscv.MSTATUS.SetBits(riscv.MSTATUS_MIE)
// Wait for the interrupt to happen.
for {
riscv.Asm("wfi")
}
}
// We got awoken.
// Lock again, after we finished sleeping.
schedulerLock.Lock()
sleepingCore = 0xff
}
// Interrupt an ongoing call to sleepTicksMulticore on another core.
// This may only be called with the scheduler lock held.
func interruptSleepTicksMulticore(wakeup timeUnit) {
if sleepingCore != 0xff {
// Immediately exit the sleep.
old := aclintMTIMECMP[sleepingCore].Get()
if uint64(wakeup) < old {
aclintMTIMECMP[sleepingCore].Set(uint64(wakeup))
}
riscv.Asm("wfi")
}
}
@@ -98,7 +303,7 @@ func ticks() timeUnit {
return timeUnit(lowBits) | (timeUnit(highBits) << 32)
}
// Retry, because there was a rollover in the low bits (happening every
// 429 days).
// ~7 days).
highBits = newHighBits
}
}
@@ -120,7 +325,10 @@ var (
low volatile.Register32
high volatile.Register32
})(unsafe.Pointer(uintptr(0x0200_bff8)))
aclintMTIMECMP = (*volatile.Register64)(unsafe.Pointer(uintptr(0x0200_4000)))
aclintMTIMECMP = (*[4095]volatile.Register64)(unsafe.Pointer(uintptr(0x0200_4000)))
aclintMSWI = (*struct {
MSIP [4095]volatile.Register32
})(unsafe.Pointer(uintptr(0x0200_0000)))
)
func putchar(c byte) {
@@ -137,17 +345,166 @@ func buffered() int {
return 0
}
// Define the various spinlocks needed by the runtime.
var (
schedulerLock spinLock
futexLock spinLock
atomicsLock spinLock
printLock spinLock
)
type spinLock struct {
atomic.Uint32
}
func (l *spinLock) Lock() {
// Try to replace 0 with 1. Once we succeed, the lock has been acquired.
for !l.Uint32.CompareAndSwap(0, 1) {
spinLoopHint()
}
}
func (l *spinLock) Unlock() {
// Safety check: the spinlock should have been locked.
if schedulerAsserts && l.Uint32.Load() != 1 {
runtimePanic("unlock of unlocked spinlock")
}
// Unlock the lock. Simply write 0, because we already know it is locked.
l.Uint32.Store(0)
}
// Hint to the CPU that this core is just waiting, and the core can go into a
// lower energy state.
func spinLoopHint() {
// This is a no-op in QEMU TCG (but added here for completeness):
// https://github.com/qemu/qemu/blob/v9.2.3/target/riscv/insn_trans/trans_rvi.c.inc#L856
riscv.Asm("pause")
}
func currentCPU() uint32 {
return uint32(riscv.MHARTID.Get())
}
func startSecondaryCores() {
// Start all the other cores besides hart 0.
for hart := 1; hart < numCPU; hart++ {
// Signal the given hart it is ready to start using a software
// interrupt.
aclintMSWI.MSIP[hart].Set(1)
}
}
// Bitset of harts that are currently sleeping in schedulerUnlockAndWait.
// This supports up to 8 harts.
// This variable may only be accessed with the scheduler lock held.
var sleepingHarts uint8
// Checkpoints for cores waiting for runnable tasks.
var schedulerWaitCheckpoints [numCPU]interrupt.Checkpoint
// Put the scheduler to sleep, since there are no tasks to run.
// This will unlock the scheduler lock, and must be called with the scheduler
// lock held.
func schedulerUnlockAndWait() {
hartID := currentCPU()
// Mark the current hart as sleeping.
sleepingHarts |= uint8(1 << hartID)
// If this is the last core awake and is going to sleep, the scheduler is
// deadlocked.
// We can do this check since this is not baremetal: there won't be any
// external interrupts that might unblock a goroutine.
if sleepingHarts == (1<<numCPU)-1 {
runtimePanic("all cores are sleeping - deadlock!")
}
// Need to disable interrupts while saving the checkpoint, otherwise if the
// software interrupt happens earlier for another reason (e.g. a GC cycle)
// it will see an incomplete checkpoint and the schedulerLock might not be
// unlocked yet. That will lead to an invalid state.
riscv.MSTATUS.ClearBits(riscv.MSTATUS_MIE)
if schedulerWaitCheckpoints[hartID].Save() {
schedulerLock.Unlock()
riscv.MSTATUS.SetBits(riscv.MSTATUS_MIE)
// Wait until we get awoken :)
for {
riscv.Asm("wfi")
}
}
// We got awoken again. We need to lock the scheduler again before
// returning.
schedulerLock.Lock()
}
// Wake another core, if one is sleeping. Must be called with the scheduler lock
// held.
func schedulerWake() {
// Look up the lowest-numbered hart that is sleeping.
// Returns 8 if there are no sleeping harts.
hart := bits.TrailingZeros8(sleepingHarts)
if hart < 8 {
// There is a sleeping hart. Wake it.
sleepingHarts &^= 1 << hart // clear the bit
aclintMSWI.MSIP[hart].Set(1) // send software interrupt
}
}
// Pause the given core by sending it an interrupt.
func gcPauseCore(core uint32) {
aclintMSWI.MSIP[core].Set(1) // send software interrupt
}
// Signal the given core that it can resume one step.
// This is called twice after gcPauseCore: the first time to scan the stack of
// the core, and the second time to end the stop-the-world phase.
func gcSignalCore(core uint32) {
aclintMSWI.MSIP[core].Set(1) // send software interrupt
}
func abort() {
exit(1)
}
// Zero in the default state, when non-zero it indicates the exit code plus one.
// So exit(0) will result in 1, exit(1) in 2, etc.
var exitCodePlusOne atomic.Uint32
func exit(code int) {
// Check for invalid values, to be sure.
if code < 0 {
code = 255
}
// If we're not on hart 0, we can't exit QEMU.
// Therefore, send an interrupt to hart 0 instead to request an exit.
if currentCPU() != 0 {
// Signal hart 0 to exit.
exitCodePlusOne.Store(uint32(code) + 1)
aclintMSWI.MSIP[0].Set(1)
// Wait for the interrupt to happen. This should happen immediately.
for {
riscv.Asm("wfi")
}
}
exitNow(uint32(code))
}
// Send an exit signal to the test finisher pseudo-device, without checking
// whether we are on hart 0.
func exitNow(code uint32) {
// Make sure the QEMU process exits.
if code == 0 {
testFinisher.Set(0x5555) // FINISHER_PASS
} else {
// Exit code is stored in the upper 16 bits of the 32 bit value.
testFinisher.Set(uint32(code)<<16 | 0x3333) // FINISHER_FAIL
testFinisher.Set(code<<16 | 0x3333) // FINISHER_FAIL
}
// Lock up forever (as a fallback).
@@ -162,10 +519,6 @@ func exit(code int) {
func handleException(code uint) {
// For a list of exception codes, see:
// https://content.riscv.org/wp-content/uploads/2019/08/riscv-privileged-20190608-1.pdf#page=49
print("fatal error: exception with mcause=")
print(code)
print(" pc=")
print(riscv.MEPC.Get())
println()
print("fatal error: exception with mcause=", code, " pc=", riscv.MEPC.Get(), " hart=", uint(riscv.MHARTID.Get()), "\r\n")
abort()
}
+16
View File
@@ -252,3 +252,19 @@ func run() {
}()
scheduler(false)
}
func lockAtomics() interrupt.State {
return interrupt.Disable()
}
func unlockAtomics(mask interrupt.State) {
interrupt.Restore(mask)
}
func printlock() {
// nothing to do
}
func printunlock() {
// nothing to do
}
+317
View File
@@ -0,0 +1,317 @@
//go:build scheduler.cores
package runtime
import (
"internal/task"
"runtime/interrupt"
"sync/atomic"
)
const hasScheduler = true
const hasParallelism = true
var mainExited atomic.Uint32
// Which task is running on a given core (or nil if there is no task running on
// the core).
var cpuTasks [numCPU]*task.Task
var (
sleepQueue *task.Task
runqueue task.Queue
)
func deadlock() {
// Call yield without requesting a wakeup.
task.Pause()
trap()
}
// Mark the given task as ready to resume.
// This is allowed even if the task isn't paused yet, but will pause soon.
func scheduleTask(t *task.Task) {
schedulerLock.Lock()
switch t.RunState {
case task.RunStatePaused:
// Paused, state is saved on the stack.
// Add it to the runqueue...
runqueue.Push(t)
// ...and wake up a sleeping core, if there is one.
// (If all cores are already busy, this is a no-op).
schedulerWake()
case task.RunStateRunning:
// Not yet paused (probably going to pause very soon), so let the
// Pause() function know it can resume immediately.
t.RunState = task.RunStateResuming
default:
if schedulerAsserts {
runtimePanic("scheduler: unknown run state")
}
}
schedulerLock.Unlock()
}
func addSleepTask(t *task.Task, wakeup timeUnit) {
// Save the timestamp when the task should be woken up.
t.Data = uint64(wakeup)
// If another core is currently using the timer, make sure it wakes up at
// the right time.
interruptSleepTicksMulticore(wakeup)
// Find the position where we should insert this task in the queue.
q := &sleepQueue
for {
if *q == nil {
// Found the end of the time queue. Insert it here, at the end.
break
}
if timeUnit((*q).Data) > timeUnit(t.Data) {
// Found a task in the queue that has a timeout before the
// to-be-sleeping task. Insert our task right before.
break
}
q = &(*q).Next
}
// Insert the task into the queue (this could be at the end, if *q is nil).
t.Next = *q
*q = t
}
func Gosched() {
schedulerLock.Lock()
runqueue.Push(task.Current())
task.PauseLocked()
}
func addTimer(tn *timerNode) {
schedulerLock.Lock()
timerQueueAdd(tn)
interruptSleepTicksMulticore(tn.whenTicks())
schedulerLock.Unlock()
}
func removeTimer(t *timer) *timerNode {
schedulerLock.Lock()
n := timerQueueRemove(t)
schedulerLock.Unlock()
return n
}
func schedulerRunQueue() *task.Queue {
return &runqueue
}
// Pause the current task for a given time.
//
//go:linkname sleep time.Sleep
func sleep(duration int64) {
if duration <= 0 {
return
}
wakeup := ticks() + nanosecondsToTicks(duration)
// While the scheduler is locked:
// - add this task to the sleep queue
// - switch to the scheduler (only allowed while locked)
// - let the scheduler handle it from there
schedulerLock.Lock()
addSleepTask(task.Current(), wakeup)
task.PauseLocked()
}
// This function is called on the first core in the system. It will wake up the
// other cores when ready.
func run() {
initHeap()
go func() {
// Package initializers are currently run single-threaded.
// This might help with registering interrupts and such.
initAll()
// After package initializers have finished, start all the other cores.
startSecondaryCores()
// Run main.main.
callMain()
// main.main has exited, so the program should exit.
mainExited.Store(1)
}()
// The scheduler must always be entered while the scheduler lock is taken.
schedulerLock.Lock()
scheduler(false)
schedulerLock.Unlock()
}
func scheduler(_ bool) {
for mainExited.Load() == 0 {
// Check for ready-to-run tasks.
if runnable := runqueue.Pop(); runnable != nil {
// Resume it now.
setCurrentTask(runnable)
runnable.RunState = task.RunStateRunning
schedulerLock.Unlock() // unlock before resuming, Pause() will lock again
runnable.Resume()
setCurrentTask(nil)
continue
}
var now timeUnit
if sleepQueue != nil || timerQueue != nil {
now = ticks()
// Check whether the first task in the sleep queue is ready to run.
if sleepingTask := sleepQueue; sleepingTask != nil && now >= timeUnit(sleepingTask.Data) {
// It is, pop it from the queue.
sleepQueue = sleepQueue.Next
sleepingTask.Next = nil
// Run it now.
setCurrentTask(sleepingTask)
sleepingTask.RunState = task.RunStateRunning
schedulerLock.Unlock() // unlock before resuming, Pause() will lock again
sleepingTask.Resume()
setCurrentTask(nil)
continue
}
// Check whether a timer has expired that needs to be run.
if timerQueue != nil && now >= timerQueue.whenTicks() {
delay := ticksToNanoseconds(now - timerQueue.whenTicks())
// Pop timer from queue.
tn := timerQueue
timerQueue = tn.next
tn.next = nil
// Run the callback stored in this timer node.
schedulerLock.Unlock()
tn.callback(tn, delay)
schedulerLock.Lock()
continue
}
}
// At this point, there are no runnable tasks anymore.
// If another core is using the clock, let it handle the sleep queue.
if hasSleepingCore() {
schedulerUnlockAndWait()
continue
}
// The timer is free to use, so check whether there are any future
// tasks/timers that we can wait for.
var timeLeft timeUnit
if sleepingTask := sleepQueue; sleepingTask != nil {
// We already checked that there is no ready-to-run sleeping task
// (using the same 'now' value), so timeLeft will always be
// positive.
timeLeft = timeUnit(sleepingTask.Data) - now
}
if timerQueue != nil {
// If the timer queue needs to run earlier, reduce the time we are
// going to sleep.
// Like with sleepQueue, we already know there is no timer ready to
// run since we already checked above.
timeLeftForTimer := timerQueue.whenTicks() - now
if sleepQueue == nil || timeLeftForTimer < timeLeft {
timeLeft = timeLeftForTimer
}
}
if timeLeft > 0 {
// Sleep for a bit until the next task or timer is ready to run.
sleepTicksMulticore(timeLeft)
continue
}
// No runnable tasks and no sleeping tasks or timers. There's nothing to
// do.
// Wait until something happens (like an interrupt).
schedulerUnlockAndWait()
}
}
func currentTask() *task.Task {
return cpuTasks[currentCPU()]
}
func setCurrentTask(task *task.Task) {
cpuTasks[currentCPU()] = task
}
func lockScheduler() {
schedulerLock.Lock()
}
func unlockScheduler() {
schedulerLock.Unlock()
}
func lockFutex() interrupt.State {
mask := interrupt.Disable()
futexLock.Lock()
return mask
}
func unlockFutex(state interrupt.State) {
futexLock.Unlock()
interrupt.Restore(state)
}
// Use a single spinlock for atomics. This works fine, since atomics are very
// short sequences of instructions.
func lockAtomics() interrupt.State {
mask := interrupt.Disable()
atomicsLock.Lock()
return mask
}
func unlockAtomics(mask interrupt.State) {
atomicsLock.Unlock()
interrupt.Restore(mask)
}
var systemStack [numCPU]uintptr
// Implementation detail of the internal/task package.
// It needs to store the system stack pointer somewhere, and needs to know how
// many cores there are to do so. But it doesn't know the number of cores. Hence
// why this is implemented in the runtime.
func systemStackPtr() *uintptr {
return &systemStack[currentCPU()]
}
// Color the 'print' and 'println' output according to the current CPU.
// This may be helpful for debugging, but should be disabled otherwise.
const cpuColoredPrint = false
func printlock() {
printLock.Lock()
if cpuColoredPrint {
switch currentCPU() {
case 1:
printstring("\x1b[32m") // green
case 2:
printstring("\x1b[33m") // yellow
case 3:
printstring("\x1b[34m") // blue
}
}
}
func printunlock() {
if cpuColoredPrint {
if currentCPU() != 0 {
printstring("\x1b[0m") // reset colored output
}
}
printLock.Unlock()
}
+20 -1
View File
@@ -2,7 +2,10 @@
package runtime
import "internal/task"
import (
"internal/task"
"runtime/interrupt"
)
const hasScheduler = false
@@ -73,3 +76,19 @@ func scheduler(returnAtDeadlock bool) {
// this code should be unreachable.
runtimePanic("unreachable: scheduler must not be called with the 'none' scheduler")
}
func lockAtomics() interrupt.State {
return interrupt.Disable()
}
func unlockAtomics(mask interrupt.State) {
interrupt.Restore(mask)
}
func printlock() {
// nothing to do
}
func printunlock() {
// nothing to do
}
+13
View File
@@ -0,0 +1,13 @@
//go:build scheduler.tasks
package runtime
var systemStack uintptr
// Implementation detail of the internal/task package.
// It needs to store the system stack pointer somewhere, and needs to know how
// many cores there are to do so. But it doesn't know the number of cores. Hence
// why this is implemented in the runtime.
func systemStackPtr() *uintptr {
return &systemStack
}
+31 -1
View File
@@ -2,7 +2,10 @@
package runtime
import "internal/task"
import (
"internal/task"
"runtime/interrupt"
)
const hasScheduler = false // not using the cooperative scheduler
@@ -127,3 +130,30 @@ func runqueueForGC() *task.Queue {
// There is only a runqueue when using the cooperative scheduler.
return nil
}
// Lock to make sure print calls do not interleave.
var printLock task.Mutex
func printlock() {
printLock.Lock()
}
func printunlock() {
printLock.Unlock()
}
// The atomics lock isn't used as a lock for actual atomics. It is used inside
// internal/task.Stack and internal/task.Queue to make sure their operations are
// actually atomic. (This might not actually be needed, since the use in
// sync.Cond doesn't need atomicity).
var atomicsLock task.Mutex
func lockAtomics() interrupt.State {
atomicsLock.Lock()
return 0
}
func unlockAtomics(mask interrupt.State) {
atomicsLock.Unlock()
}
+17 -4
View File
@@ -1,8 +1,21 @@
{
"inherits": ["riscv32"],
"features": "+32bit,+a,+c,+m,+zmmul,-b,-d,-e,-experimental-smmpm,-experimental-smnpm,-experimental-ssnpm,-experimental-sspm,-experimental-ssqosid,-experimental-supm,-experimental-zacas,-experimental-zalasr,-experimental-zicfilp,-experimental-zicfiss,-f,-h,-relax,-shcounterenw,-shgatpa,-shtvala,-shvsatpa,-shvstvala,-shvstvecd,-smaia,-smcdeleg,-smcsrind,-smepmp,-smstateen,-ssaia,-ssccfg,-ssccptr,-sscofpmf,-sscounterenw,-sscsrind,-ssstateen,-ssstrict,-sstc,-sstvala,-sstvecd,-ssu64xl,-svade,-svadu,-svbare,-svinval,-svnapot,-svpbmt,-v,-xcvalu,-xcvbi,-xcvbitmanip,-xcvelw,-xcvmac,-xcvmem,-xcvsimd,-xesppie,-xsfcease,-xsfvcp,-xsfvfnrclipxfqf,-xsfvfwmaccqqq,-xsfvqmaccdod,-xsfvqmaccqoq,-xsifivecdiscarddlone,-xsifivecflushdlone,-xtheadba,-xtheadbb,-xtheadbs,-xtheadcmo,-xtheadcondmov,-xtheadfmemidx,-xtheadmac,-xtheadmemidx,-xtheadmempair,-xtheadsync,-xtheadvdot,-xventanacondops,-xwchc,-za128rs,-za64rs,-zaamo,-zabha,-zalrsc,-zama16b,-zawrs,-zba,-zbb,-zbc,-zbkb,-zbkc,-zbkx,-zbs,-zca,-zcb,-zcd,-zce,-zcf,-zcmop,-zcmp,-zcmt,-zdinx,-zfa,-zfbfmin,-zfh,-zfhmin,-zfinx,-zhinx,-zhinxmin,-zic64b,-zicbom,-zicbop,-zicboz,-ziccamoa,-ziccif,-zicclsm,-ziccrse,-zicntr,-zicond,-zicsr,-zifencei,-zihintntl,-zihintpause,-zihpm,-zimop,-zk,-zkn,-zknd,-zkne,-zknh,-zkr,-zks,-zksed,-zksh,-zkt,-ztso,-zvbb,-zvbc,-zve32f,-zve32x,-zve64d,-zve64f,-zve64x,-zvfbfmin,-zvfbfwma,-zvfh,-zvfhmin,-zvkb,-zvkg,-zvkn,-zvknc,-zvkned,-zvkng,-zvknha,-zvknhb,-zvks,-zvksc,-zvksed,-zvksg,-zvksh,-zvkt,-zvl1024b,-zvl128b,-zvl16384b,-zvl2048b,-zvl256b,-zvl32768b,-zvl32b,-zvl4096b,-zvl512b,-zvl64b,-zvl65536b,-zvl8192b",
"build-tags": ["virt", "qemu"],
"inherits": [
"riscv32"
],
"features": "+32bit,+a,+c,+m,+zihintpause,+zmmul,-b,-d,-e,-experimental-smmpm,-experimental-smnpm,-experimental-ssnpm,-experimental-sspm,-experimental-ssqosid,-experimental-supm,-experimental-zacas,-experimental-zalasr,-experimental-zicfilp,-experimental-zicfiss,-f,-h,-relax,-shcounterenw,-shgatpa,-shtvala,-shvsatpa,-shvstvala,-shvstvecd,-smaia,-smcdeleg,-smcsrind,-smepmp,-smstateen,-ssaia,-ssccfg,-ssccptr,-sscofpmf,-sscounterenw,-sscsrind,-ssstateen,-ssstrict,-sstc,-sstvala,-sstvecd,-ssu64xl,-svade,-svadu,-svbare,-svinval,-svnapot,-svpbmt,-v,-xcvalu,-xcvbi,-xcvbitmanip,-xcvelw,-xcvmac,-xcvmem,-xcvsimd,-xesppie,-xsfcease,-xsfvcp,-xsfvfnrclipxfqf,-xsfvfwmaccqqq,-xsfvqmaccdod,-xsfvqmaccqoq,-xsifivecdiscarddlone,-xsifivecflushdlone,-xtheadba,-xtheadbb,-xtheadbs,-xtheadcmo,-xtheadcondmov,-xtheadfmemidx,-xtheadmac,-xtheadmemidx,-xtheadmempair,-xtheadsync,-xtheadvdot,-xventanacondops,-xwchc,-za128rs,-za64rs,-zaamo,-zabha,-zalrsc,-zama16b,-zawrs,-zba,-zbb,-zbc,-zbkb,-zbkc,-zbkx,-zbs,-zca,-zcb,-zcd,-zce,-zcf,-zcmop,-zcmp,-zcmt,-zdinx,-zfa,-zfbfmin,-zfh,-zfhmin,-zfinx,-zhinx,-zhinxmin,-zic64b,-zicbom,-zicbop,-zicboz,-ziccamoa,-ziccif,-zicclsm,-ziccrse,-zicntr,-zicond,-zicsr,-zifencei,-zihintntl,-zihpm,-zimop,-zk,-zkn,-zknd,-zkne,-zknh,-zkr,-zks,-zksed,-zksh,-zkt,-ztso,-zvbb,-zvbc,-zve32f,-zve32x,-zve64d,-zve64f,-zve64x,-zvfbfmin,-zvfbfwma,-zvfh,-zvfhmin,-zvkb,-zvkg,-zvkn,-zvknc,-zvkned,-zvkng,-zvknha,-zvknhb,-zvks,-zvksc,-zvksed,-zvksg,-zvksh,-zvkt,-zvl1024b,-zvl128b,-zvl16384b,-zvl2048b,-zvl256b,-zvl32768b,-zvl32b,-zvl4096b,-zvl512b,-zvl64b,-zvl65536b,-zvl8192b",
"build-tags": [
"virt",
"qemu"
],
"scheduler": "cores",
"default-stack-size": 8192,
"cflags": [
"-march=rv32imaczihintpause",
"-DTINYGO_CORES=4"
],
"ldflags": [
"--defsym=__num_stacks=4"
],
"linkerscript": "targets/riscv-qemu.ld",
"emulator": "qemu-system-riscv32 -machine virt,aclint=on -nographic -bios none -device virtio-rng-device -kernel {}"
"emulator": "qemu-system-riscv32 -machine virt,aclint=on -smp 4 -nographic -bios none -device virtio-rng-device -kernel {}"
}
+22 -1
View File
@@ -16,13 +16,34 @@ SECTIONS
/* Put the stack at the bottom of RAM, so that the application will
* crash on stack overflow instead of silently corrupting memory.
* See: http://blog.japaric.io/stack-overflow-protection/ */
.stack (NOLOAD) :
.stack0 (NOLOAD) :
{
. = ALIGN(16);
. += _stack_size;
_stack_top = .;
} >RAM
.stack1 (NOLOAD) :
{
. = ALIGN(16);
. += DEFINED(__num_stacks) && __num_stacks >= 2 ? _stack_size : 0;
_stack1_top = .;
} >RAM
.stack2 (NOLOAD) :
{
. = ALIGN(16);
. += DEFINED(__num_stacks) && __num_stacks >= 3 ? _stack_size : 0;
_stack2_top = .;
} >RAM
.stack3 (NOLOAD) :
{
. = ALIGN(16);
. += DEFINED(__num_stacks) && __num_stacks >= 4 ? _stack_size : 0;
_stack3_top = .;
} >RAM
/* Start address (in flash) of .data, used by startup code. */
_sidata = LOADADDR(.data);
@@ -26,7 +26,6 @@ package runtime
import (
_ "unsafe"
"runtime/interrupt"
)
// Documentation:
@@ -41,29 +40,29 @@ import (
func __atomic_load_{{.}}(ptr *uint{{$bits}}, ordering uintptr) uint{{$bits}} {
// The LLVM docs for this say that there is a val argument after the pointer.
// That is a typo, and the GCC docs omit it.
mask := interrupt.Disable()
mask := lockAtomics()
val := *ptr
interrupt.Restore(mask)
unlockAtomics(mask)
return val
}
{{end}}
{{- define "store"}}{{$bits := mul . 8 -}}
//export __atomic_store_{{.}}
func __atomic_store_{{.}}(ptr *uint{{$bits}}, val uint{{$bits}}, ordering uintptr) {
mask := interrupt.Disable()
mask := lockAtomics()
*ptr = val
interrupt.Restore(mask)
unlockAtomics(mask)
}
{{end}}
{{- define "cas"}}{{$bits := mul . 8 -}}
//go:inline
func doAtomicCAS{{$bits}}(ptr *uint{{$bits}}, expected, desired uint{{$bits}}) uint{{$bits}} {
mask := interrupt.Disable()
mask := lockAtomics()
old := *ptr
if old == expected {
*ptr = desired
}
interrupt.Restore(mask)
unlockAtomics(mask)
return old
}
@@ -82,10 +81,10 @@ func __atomic_compare_exchange_{{.}}(ptr, expected *uint{{$bits}}, desired uint{
{{- define "swap"}}{{$bits := mul . 8 -}}
//go:inline
func doAtomicSwap{{$bits}}(ptr *uint{{$bits}}, new uint{{$bits}}) uint{{$bits}} {
mask := interrupt.Disable()
mask := lockAtomics()
old := *ptr
*ptr = new
interrupt.Restore(mask)
unlockAtomics(mask)
return old
}
@@ -111,11 +110,11 @@ func __atomic_exchange_{{.}}(ptr *uint{{$bits}}, new uint{{$bits}}, ordering uin
//go:inline
func {{$opfn}}(ptr *{{$type}}, value {{$type}}) (old, new {{$type}}) {
mask := interrupt.Disable()
mask := lockAtomics()
old = *ptr
{{$opdef}}
*ptr = new
interrupt.Restore(mask)
unlockAtomics(mask)
return old, new
}