aboutsummaryrefslogtreecommitdiff
path: root/libomptarget/deviceRTLs/nvptx/src/state-queue.h
blob: 9d7576bcd76e29cb1f3af49709773341e87bc529 (plain)
1
2
3
4
5
6
7
8
9
10
11
12
13
14
15
16
17
18
19
20
21
22
23
24
25
26
27
28
29
30
31
32
33
34
35
36
37
38
39
40
41
42
43
44
45
46
47
48
49
50
51
//===--------- statequeue.h - NVPTX OpenMP GPU State Queue ------- CUDA -*-===//
//
// Part of the LLVM Project, under the Apache License v2.0 with LLVM Exceptions.
// See https://llvm.org/LICENSE.txt for license information.
// SPDX-License-Identifier: Apache-2.0 WITH LLVM-exception
//
//===----------------------------------------------------------------------===//
//
// This file contains a queue to hand out OpenMP state objects to teams of
// one or more kernels.
//
// Reference:
// Thomas R.W. Scogland and Wu-chun Feng. 2015.
// Design and Evaluation of Scalable Concurrent Queues for Many-Core
// Architectures. International Conference on Performance Engineering.
//
//===----------------------------------------------------------------------===//

#ifndef __STATE_QUEUE_H
#define __STATE_QUEUE_H

#include <stdint.h>

#include "option.h" // choices we have

template <typename ElementType, uint32_t SIZE> class omptarget_nvptx_Queue {
private:
  ElementType elements[SIZE];
  volatile ElementType *elementQueue[SIZE];
  volatile uint32_t head;
  volatile uint32_t ids[SIZE];
  volatile uint32_t tail;

  static const uint32_t MAX_ID = (1u << 31) / SIZE / 2;
  INLINE uint32_t ENQUEUE_TICKET();
  INLINE uint32_t DEQUEUE_TICKET();
  INLINE static uint32_t ID(uint32_t ticket);
  INLINE bool IsServing(uint32_t slot, uint32_t id);
  INLINE void PushElement(uint32_t slot, ElementType *element);
  INLINE ElementType *PopElement(uint32_t slot);
  INLINE void DoneServing(uint32_t slot, uint32_t id);

public:
  INLINE omptarget_nvptx_Queue() {}
  INLINE void Enqueue(ElementType *element);
  INLINE ElementType *Dequeue();
};

#include "state-queuei.h"

#endif