forked from OSchip/llvm-project
AMDGPU: fix missing s_waitcnt
Summary: The pass that inserts s_waitcnt instructions where needed propagated info used to track dependencies for each block by iterating over the predecessor blocks. The iteration was terminated when a predecessor that had not yet been processed was encountered. Any info in blocks later in the list was therefore not processed, leading to the possiblility of a required s_waitcnt not being inserted. The fix is simply to change the "break" to "continue" for the relevant loops, so that all visited blocks are processed. This is likely what was intended when the code was written. There is no test case provided for this fix because: 1) the only example that reproduces this is large and resistant to being reduced 2) the change is trivial Subscribers: arsenm, kzhuravl, wdng, nhaehnle, yaxunl, dstuttard, tpr, t-tye Differential Revision: https://reviews.llvm.org/D40544 llvm-svn: 319651
This commit is contained in:
parent
64fc9cf2e5
commit
6c6d5e24cd
|
@ -1269,7 +1269,7 @@ void SIInsertWaitcnts::mergeInputScoreBrackets(MachineBasicBlock &Block) {
|
|||
BlockWaitcntBracketsMap[pred].get();
|
||||
bool Visited = BlockVisitedSet.find(pred) != BlockVisitedSet.end();
|
||||
if (!Visited || PredScoreBrackets->getWaitAtBeginning()) {
|
||||
break;
|
||||
continue;
|
||||
}
|
||||
for (enum InstCounterType T = VM_CNT; T < NUM_INST_CNTS;
|
||||
T = (enum InstCounterType)(T + 1)) {
|
||||
|
@ -1308,7 +1308,7 @@ void SIInsertWaitcnts::mergeInputScoreBrackets(MachineBasicBlock &Block) {
|
|||
BlockWaitcntBracketsMap[Pred].get();
|
||||
bool Visited = BlockVisitedSet.find(Pred) != BlockVisitedSet.end();
|
||||
if (!Visited || PredScoreBrackets->getWaitAtBeginning()) {
|
||||
break;
|
||||
continue;
|
||||
}
|
||||
|
||||
int GDSSpan = PredScoreBrackets->getEventUB(GDS_GPR_LOCK) -
|
||||
|
@ -1355,7 +1355,7 @@ void SIInsertWaitcnts::mergeInputScoreBrackets(MachineBasicBlock &Block) {
|
|||
// Set the register scoreboard.
|
||||
for (MachineBasicBlock *Pred : Block.predecessors()) {
|
||||
if (BlockVisitedSet.find(Pred) == BlockVisitedSet.end()) {
|
||||
break;
|
||||
continue;
|
||||
}
|
||||
|
||||
BlockWaitcntBrackets *PredScoreBrackets =
|
||||
|
@ -1469,7 +1469,7 @@ void SIInsertWaitcnts::mergeInputScoreBrackets(MachineBasicBlock &Block) {
|
|||
// the delayed nature of these operations.
|
||||
for (MachineBasicBlock *Pred : Block.predecessors()) {
|
||||
if (BlockVisitedSet.find(Pred) == BlockVisitedSet.end()) {
|
||||
break;
|
||||
continue;
|
||||
}
|
||||
|
||||
BlockWaitcntBrackets *PredScoreBrackets =
|
||||
|
|
Loading…
Reference in New Issue