2016-07-27 00:45:58 +08:00
|
|
|
//===-- AMDGPUMachineFunctionInfo.cpp ---------------------------------------=//
|
|
|
|
//
|
2019-01-19 16:50:56 +08:00
|
|
|
// Part of the LLVM Project, under the Apache License v2.0 with LLVM Exceptions.
|
|
|
|
// See https://llvm.org/LICENSE.txt for license information.
|
|
|
|
// SPDX-License-Identifier: Apache-2.0 WITH LLVM-exception
|
2016-07-27 00:45:58 +08:00
|
|
|
//
|
|
|
|
//===----------------------------------------------------------------------===//
|
|
|
|
|
2013-04-02 05:47:53 +08:00
|
|
|
#include "AMDGPUMachineFunction.h"
|
2016-07-27 00:45:58 +08:00
|
|
|
#include "AMDGPUSubtarget.h"
|
2018-05-26 01:25:12 +08:00
|
|
|
#include "AMDGPUPerfHintAnalysis.h"
|
|
|
|
#include "llvm/CodeGen/MachineModuleInfo.h"
|
2016-06-18 13:15:53 +08:00
|
|
|
|
2013-07-17 08:31:35 +08:00
|
|
|
using namespace llvm;
|
2013-04-02 05:47:53 +08:00
|
|
|
|
|
|
|
AMDGPUMachineFunction::AMDGPUMachineFunction(const MachineFunction &MF) :
|
2014-07-13 11:06:39 +08:00
|
|
|
MachineFunctionInfo(),
|
2019-11-18 19:18:07 +08:00
|
|
|
Mode(MF.getFunction()),
|
2017-12-16 06:22:58 +08:00
|
|
|
IsEntryFunction(AMDGPU::isEntryFunctionCC(MF.getFunction().getCallingConv())),
|
2020-05-20 01:44:14 +08:00
|
|
|
NoSignedZerosFPMath(MF.getTarget().Options.NoSignedZerosFPMath) {
|
2018-07-20 17:05:08 +08:00
|
|
|
const AMDGPUSubtarget &ST = AMDGPUSubtarget::get(MF);
|
|
|
|
|
2016-07-27 00:45:58 +08:00
|
|
|
// FIXME: Should initialize KernArgSize based on ExplicitKernelArgOffset,
|
|
|
|
// except reserved size is not correctly aligned.
|
2018-07-20 17:05:08 +08:00
|
|
|
const Function &F = MF.getFunction();
|
2018-05-26 01:25:12 +08:00
|
|
|
|
2019-07-06 04:26:13 +08:00
|
|
|
Attribute MemBoundAttr = F.getFnAttribute("amdgpu-memory-bound");
|
|
|
|
MemoryBound = MemBoundAttr.isStringAttribute() &&
|
|
|
|
MemBoundAttr.getValueAsString() == "true";
|
|
|
|
|
|
|
|
Attribute WaveLimitAttr = F.getFnAttribute("amdgpu-wave-limiter");
|
|
|
|
WaveLimiter = WaveLimitAttr.isStringAttribute() &&
|
|
|
|
WaveLimitAttr.getValueAsString() == "true";
|
2018-07-20 17:05:08 +08:00
|
|
|
|
|
|
|
CallingConv::ID CC = F.getCallingConv();
|
|
|
|
if (CC == CallingConv::AMDGPU_KERNEL || CC == CallingConv::SPIR_KERNEL)
|
|
|
|
ExplicitKernArgSize = ST.getExplicitKernArgSize(F, MaxKernArgAlign);
|
2016-07-01 18:00:58 +08:00
|
|
|
}
|
|
|
|
|
2016-07-27 00:45:58 +08:00
|
|
|
unsigned AMDGPUMachineFunction::allocateLDSGlobal(const DataLayout &DL,
|
2020-05-19 11:38:13 +08:00
|
|
|
const GlobalVariable &GV) {
|
2016-07-27 00:45:58 +08:00
|
|
|
auto Entry = LocalMemoryObjects.insert(std::make_pair(&GV, 0));
|
|
|
|
if (!Entry.second)
|
|
|
|
return Entry.first->second;
|
|
|
|
|
2020-06-29 19:56:06 +08:00
|
|
|
Align Alignment =
|
|
|
|
DL.getValueOrABITypeAlignment(GV.getAlign(), GV.getValueType());
|
2016-07-27 00:45:58 +08:00
|
|
|
|
|
|
|
/// TODO: We should sort these to minimize wasted space due to alignment
|
|
|
|
/// padding. Currently the padding is decided by the first encountered use
|
|
|
|
/// during lowering.
|
[amdgpu] Add codegen support for HIP dynamic shared memory.
Summary:
- HIP uses an unsized extern array `extern __shared__ T s[]` to declare
the dynamic shared memory, which size is not known at the
compile time.
Reviewers: arsenm, yaxunl, kpyzhov, b-sumner
Subscribers: kzhuravl, jvesely, wdng, nhaehnle, dstuttard, tpr, t-tye, hiraditya, kerbowa, llvm-commits
Tags: #llvm
Differential Revision: https://reviews.llvm.org/D82496
2020-06-25 00:13:10 +08:00
|
|
|
unsigned Offset = StaticLDSSize = alignTo(StaticLDSSize, Alignment);
|
2016-07-27 00:45:58 +08:00
|
|
|
|
|
|
|
Entry.first->second = Offset;
|
[amdgpu] Add codegen support for HIP dynamic shared memory.
Summary:
- HIP uses an unsized extern array `extern __shared__ T s[]` to declare
the dynamic shared memory, which size is not known at the
compile time.
Reviewers: arsenm, yaxunl, kpyzhov, b-sumner
Subscribers: kzhuravl, jvesely, wdng, nhaehnle, dstuttard, tpr, t-tye, hiraditya, kerbowa, llvm-commits
Tags: #llvm
Differential Revision: https://reviews.llvm.org/D82496
2020-06-25 00:13:10 +08:00
|
|
|
StaticLDSSize += DL.getTypeAllocSize(GV.getValueType());
|
|
|
|
|
|
|
|
// Update the LDS size considering the padding to align the dynamic shared
|
|
|
|
// memory.
|
|
|
|
LDSSize = alignTo(StaticLDSSize, DynLDSAlign);
|
2016-07-27 00:45:58 +08:00
|
|
|
|
|
|
|
return Offset;
|
2016-07-01 18:00:58 +08:00
|
|
|
}
|
[amdgpu] Add codegen support for HIP dynamic shared memory.
Summary:
- HIP uses an unsized extern array `extern __shared__ T s[]` to declare
the dynamic shared memory, which size is not known at the
compile time.
Reviewers: arsenm, yaxunl, kpyzhov, b-sumner
Subscribers: kzhuravl, jvesely, wdng, nhaehnle, dstuttard, tpr, t-tye, hiraditya, kerbowa, llvm-commits
Tags: #llvm
Differential Revision: https://reviews.llvm.org/D82496
2020-06-25 00:13:10 +08:00
|
|
|
|
|
|
|
void AMDGPUMachineFunction::setDynLDSAlign(const DataLayout &DL,
|
|
|
|
const GlobalVariable &GV) {
|
|
|
|
assert(DL.getTypeAllocSize(GV.getValueType()).isZero());
|
|
|
|
|
|
|
|
Align Alignment =
|
|
|
|
DL.getValueOrABITypeAlignment(GV.getAlign(), GV.getValueType());
|
|
|
|
if (Alignment <= DynLDSAlign)
|
|
|
|
return;
|
|
|
|
|
|
|
|
LDSSize = alignTo(StaticLDSSize, Alignment);
|
|
|
|
DynLDSAlign = Alignment;
|
|
|
|
}
|