mirror of
https://github.com/intel/llvm.git
synced 2026-01-14 11:57:39 +08:00
[BOLT][NFC] Move DynoStats out of BinaryFunction
Summary: Move DynoStats into separate source files. (cherry picked from FBD15138883)
This commit is contained in:
@@ -13,6 +13,7 @@
|
||||
#include "BinaryBasicBlock.h"
|
||||
#include "BinaryFunction.h"
|
||||
#include "DataReader.h"
|
||||
#include "DynoStats.h"
|
||||
#include "MCPlusBuilder.h"
|
||||
#include "llvm/ADT/edit_distance.h"
|
||||
#include "llvm/ADT/StringRef.h"
|
||||
@@ -90,14 +91,6 @@ DotToolTipCode("dot-tooltip-code",
|
||||
cl::Hidden,
|
||||
cl::cat(BoltCategory));
|
||||
|
||||
static cl::opt<uint32_t>
|
||||
DynoStatsScale("dyno-stats-scale",
|
||||
cl::desc("scale to be applied while reporting dyno stats"),
|
||||
cl::Optional,
|
||||
cl::init(1),
|
||||
cl::Hidden,
|
||||
cl::cat(BoltCategory));
|
||||
|
||||
cl::opt<JumpTableSupportLevel>
|
||||
JumpTables("jump-tables",
|
||||
cl::desc("jump tables support (default=basic)"),
|
||||
@@ -193,7 +186,6 @@ bool shouldPrint(const BinaryFunction &Function) {
|
||||
namespace llvm {
|
||||
namespace bolt {
|
||||
|
||||
constexpr const char *DynoStats::Desc[];
|
||||
constexpr unsigned BinaryFunction::MinAlign;
|
||||
const char BinaryFunction::TimerGroupName[] = "buildfuncs";
|
||||
const char BinaryFunction::TimerGroupDesc[] = "Build Binary Functions";
|
||||
@@ -245,31 +237,6 @@ SMLoc findDebugLineInformationForInstructionAt(
|
||||
|
||||
} // namespace
|
||||
|
||||
bool DynoStats::operator<(const DynoStats &Other) const {
|
||||
return std::lexicographical_compare(
|
||||
&Stats[FIRST_DYNO_STAT], &Stats[LAST_DYNO_STAT],
|
||||
&Other.Stats[FIRST_DYNO_STAT], &Other.Stats[LAST_DYNO_STAT]
|
||||
);
|
||||
}
|
||||
|
||||
bool DynoStats::operator==(const DynoStats &Other) const {
|
||||
return std::equal(
|
||||
&Stats[FIRST_DYNO_STAT], &Stats[LAST_DYNO_STAT],
|
||||
&Other.Stats[FIRST_DYNO_STAT]
|
||||
);
|
||||
}
|
||||
|
||||
bool DynoStats::lessThan(const DynoStats &Other,
|
||||
ArrayRef<Category> Keys) const {
|
||||
return std::lexicographical_compare(
|
||||
Keys.begin(), Keys.end(),
|
||||
Keys.begin(), Keys.end(),
|
||||
[this,&Other](const Category A, const Category) {
|
||||
return Stats[A] < Other.Stats[A];
|
||||
}
|
||||
);
|
||||
}
|
||||
|
||||
uint64_t BinaryFunction::Count = 0;
|
||||
|
||||
const std::string *
|
||||
@@ -493,7 +460,7 @@ void BinaryFunction::print(raw_ostream &OS, std::string Annotation,
|
||||
|
||||
if (opts::PrintDynoStats && !BasicBlocksLayout.empty()) {
|
||||
OS << '\n';
|
||||
DynoStats dynoStats = getDynoStats();
|
||||
DynoStats dynoStats = getDynoStats(*this);
|
||||
OS << dynoStats;
|
||||
}
|
||||
|
||||
@@ -4284,145 +4251,6 @@ void BinaryFunction::printLoopInfo(raw_ostream &OS) const {
|
||||
OS << "Maximum nested loop depth: " << BLI->MaximumDepth << "\n\n";
|
||||
}
|
||||
|
||||
DynoStats BinaryFunction::getDynoStats() const {
|
||||
DynoStats Stats(/*PrintAArch64Stats*/ BC.isAArch64());
|
||||
|
||||
// Return empty-stats about the function we don't completely understand.
|
||||
if (!isSimple() || !hasValidProfile())
|
||||
return Stats;
|
||||
|
||||
// If the function was folded in non-relocation mode we keep its profile
|
||||
// for optimization. However, it should be excluded from the dyno stats.
|
||||
if (isFolded())
|
||||
return Stats;
|
||||
|
||||
// Update enumeration of basic blocks for correct detection of branch'
|
||||
// direction.
|
||||
updateLayoutIndices();
|
||||
|
||||
for (const auto &BB : layout()) {
|
||||
// The basic block execution count equals to the sum of incoming branch
|
||||
// frequencies. This may deviate from the sum of outgoing branches of the
|
||||
// basic block especially since the block may contain a function that
|
||||
// does not return or a function that throws an exception.
|
||||
const uint64_t BBExecutionCount = BB->getKnownExecutionCount();
|
||||
|
||||
// Ignore empty blocks and blocks that were not executed.
|
||||
if (BB->getNumNonPseudos() == 0 || BBExecutionCount == 0)
|
||||
continue;
|
||||
|
||||
// Count AArch64 linker-inserted veneers
|
||||
if(isAArch64Veneer())
|
||||
Stats[DynoStats::VENEER_CALLS_AARCH64] += getKnownExecutionCount();
|
||||
|
||||
// Count the number of calls by iterating through all instructions.
|
||||
for (const auto &Instr : *BB) {
|
||||
if (BC.MIB->isStore(Instr)) {
|
||||
Stats[DynoStats::STORES] += BBExecutionCount;
|
||||
}
|
||||
if (BC.MIB->isLoad(Instr)) {
|
||||
Stats[DynoStats::LOADS] += BBExecutionCount;
|
||||
}
|
||||
|
||||
if (!BC.MIB->isCall(Instr))
|
||||
continue;
|
||||
|
||||
uint64_t CallFreq = BBExecutionCount;
|
||||
if (BC.MIB->getConditionalTailCall(Instr)) {
|
||||
CallFreq =
|
||||
BC.MIB->getAnnotationWithDefault<uint64_t>(Instr, "CTCTakenCount");
|
||||
}
|
||||
Stats[DynoStats::FUNCTION_CALLS] += CallFreq;
|
||||
if (BC.MIB->isIndirectCall(Instr)) {
|
||||
Stats[DynoStats::INDIRECT_CALLS] += CallFreq;
|
||||
} else if (const auto *CallSymbol = BC.MIB->getTargetSymbol(Instr)) {
|
||||
const auto *BF = BC.getFunctionForSymbol(CallSymbol);
|
||||
if (BF && BF->isPLTFunction()) {
|
||||
Stats[DynoStats::PLT_CALLS] += CallFreq;
|
||||
|
||||
// We don't process PLT functions and hence have to adjust relevant
|
||||
// dynostats here for:
|
||||
//
|
||||
// jmp *GOT_ENTRY(%rip)
|
||||
//
|
||||
// NOTE: this is arch-specific.
|
||||
Stats[DynoStats::FUNCTION_CALLS] += CallFreq;
|
||||
Stats[DynoStats::INDIRECT_CALLS] += CallFreq;
|
||||
Stats[DynoStats::LOADS] += CallFreq;
|
||||
Stats[DynoStats::INSTRUCTIONS] += CallFreq;
|
||||
}
|
||||
}
|
||||
}
|
||||
|
||||
Stats[DynoStats::INSTRUCTIONS] += BB->getNumNonPseudos() * BBExecutionCount;
|
||||
|
||||
// Jump tables.
|
||||
const auto *LastInstr = BB->getLastNonPseudoInstr();
|
||||
if (BC.MIB->getJumpTable(*LastInstr)) {
|
||||
Stats[DynoStats::JUMP_TABLE_BRANCHES] += BBExecutionCount;
|
||||
DEBUG(
|
||||
static uint64_t MostFrequentJT;
|
||||
if (BBExecutionCount > MostFrequentJT) {
|
||||
MostFrequentJT = BBExecutionCount;
|
||||
dbgs() << "BOLT-INFO: most frequently executed jump table is in "
|
||||
<< "function " << *this << " in basic block " << BB->getName()
|
||||
<< " executed totally " << BBExecutionCount << " times.\n";
|
||||
}
|
||||
);
|
||||
continue;
|
||||
}
|
||||
|
||||
// Update stats for branches.
|
||||
const MCSymbol *TBB = nullptr;
|
||||
const MCSymbol *FBB = nullptr;
|
||||
MCInst *CondBranch = nullptr;
|
||||
MCInst *UncondBranch = nullptr;
|
||||
if (!BB->analyzeBranch(TBB, FBB, CondBranch, UncondBranch)) {
|
||||
continue;
|
||||
}
|
||||
|
||||
if (!CondBranch && !UncondBranch) {
|
||||
continue;
|
||||
}
|
||||
|
||||
// Simple unconditional branch.
|
||||
if (!CondBranch) {
|
||||
Stats[DynoStats::UNCOND_BRANCHES] += BBExecutionCount;
|
||||
continue;
|
||||
}
|
||||
|
||||
// CTCs
|
||||
if (BC.MIB->getConditionalTailCall(*CondBranch)) {
|
||||
if (BB->branch_info_begin() != BB->branch_info_end())
|
||||
Stats[DynoStats::UNCOND_BRANCHES] += BB->branch_info_begin()->Count;
|
||||
continue;
|
||||
}
|
||||
|
||||
// Conditional branch that could be followed by an unconditional branch.
|
||||
auto TakenCount = BB->getTakenBranchInfo().Count;
|
||||
if (TakenCount == COUNT_NO_PROFILE)
|
||||
TakenCount = 0;
|
||||
|
||||
auto NonTakenCount = BB->getFallthroughBranchInfo().Count;
|
||||
if (NonTakenCount == COUNT_NO_PROFILE)
|
||||
NonTakenCount = 0;
|
||||
|
||||
if (isForwardBranch(BB, BB->getConditionalSuccessor(true))) {
|
||||
Stats[DynoStats::FORWARD_COND_BRANCHES] += BBExecutionCount;
|
||||
Stats[DynoStats::FORWARD_COND_BRANCHES_TAKEN] += TakenCount;
|
||||
} else {
|
||||
Stats[DynoStats::BACKWARD_COND_BRANCHES] += BBExecutionCount;
|
||||
Stats[DynoStats::BACKWARD_COND_BRANCHES_TAKEN] += TakenCount;
|
||||
}
|
||||
|
||||
if (UncondBranch) {
|
||||
Stats[DynoStats::UNCOND_BRANCHES] += NonTakenCount;
|
||||
}
|
||||
}
|
||||
|
||||
return Stats;
|
||||
}
|
||||
|
||||
bool BinaryFunction::isAArch64Veneer() const {
|
||||
if (BasicBlocks.size() != 1)
|
||||
return false;
|
||||
@@ -4439,41 +4267,5 @@ bool BinaryFunction::isAArch64Veneer() const {
|
||||
return true;
|
||||
}
|
||||
|
||||
void DynoStats::print(raw_ostream &OS, const DynoStats *Other) const {
|
||||
auto printStatWithDelta = [&](const std::string &Name, uint64_t Stat,
|
||||
uint64_t OtherStat) {
|
||||
OS << format("%'20lld : ", Stat * opts::DynoStatsScale) << Name;
|
||||
if (Other) {
|
||||
if (Stat != OtherStat) {
|
||||
OtherStat = std::max(OtherStat, uint64_t(1)); // to prevent divide by 0
|
||||
OS << format(" (%+.1f%%)",
|
||||
( (float) Stat - (float) OtherStat ) * 100.0 /
|
||||
(float) (OtherStat) );
|
||||
} else {
|
||||
OS << " (=)";
|
||||
}
|
||||
}
|
||||
OS << '\n';
|
||||
};
|
||||
|
||||
for (auto Stat = DynoStats::FIRST_DYNO_STAT + 1;
|
||||
Stat < DynoStats::LAST_DYNO_STAT;
|
||||
++Stat) {
|
||||
|
||||
if (!PrintAArch64Stats && Stat == DynoStats::VENEER_CALLS_AARCH64)
|
||||
continue;
|
||||
|
||||
printStatWithDelta(Desc[Stat], Stats[Stat], Other ? (*Other)[Stat] : 0);
|
||||
}
|
||||
}
|
||||
|
||||
void DynoStats::operator+=(const DynoStats &Other) {
|
||||
for (auto Stat = DynoStats::FIRST_DYNO_STAT + 1;
|
||||
Stat < DynoStats::LAST_DYNO_STAT;
|
||||
++Stat) {
|
||||
Stats[Stat] += Other[Stat];
|
||||
}
|
||||
}
|
||||
|
||||
} // namespace bolt
|
||||
} // namespace llvm
|
||||
|
||||
Reference in New Issue
Block a user