Skip to content
体验新版
项目
组织
正在加载...
登录
切换导航
打开侧边栏
机器未来
Paddle
提交
ab049978
P
Paddle
项目概览
机器未来
/
Paddle
与 Fork 源项目一致
Fork自
PaddlePaddle / Paddle
通知
1
Star
1
Fork
0
代码
文件
提交
分支
Tags
贡献者
分支图
Diff
Issue
1
列表
看板
标记
里程碑
合并请求
0
Wiki
0
Wiki
分析
仓库
DevOps
项目成员
Pages
P
Paddle
项目概览
项目概览
详情
发布
仓库
仓库
文件
提交
分支
标签
贡献者
分支图
比较
Issue
1
Issue
1
列表
看板
标记
里程碑
合并请求
0
合并请求
0
Pages
分析
分析
仓库分析
DevOps
Wiki
0
Wiki
成员
成员
收起侧边栏
关闭侧边栏
动态
分支图
创建新Issue
提交
Issue看板
未验证
提交
ab049978
编写于
1月 05, 2021
作者:
W
WangXi
提交者:
GitHub
1月 05, 2021
浏览文件
操作
浏览文件
下载
电子邮件补丁
差异文件
[fleet] combine amp and gradient merge, test=develop (#30086)
上级
88e6dc4a
变更
4
显示空白变更内容
内联
并排
Showing
4 changed file
with
18 addition
and
4 deletion
+18
-4
python/paddle/distributed/fleet/meta_optimizers/amp_optimizer.py
...paddle/distributed/fleet/meta_optimizers/amp_optimizer.py
+0
-1
python/paddle/distributed/fleet/meta_optimizers/gradient_merge_optimizer.py
...ributed/fleet/meta_optimizers/gradient_merge_optimizer.py
+1
-0
python/paddle/fluid/contrib/mixed_precision/decorator.py
python/paddle/fluid/contrib/mixed_precision/decorator.py
+4
-3
python/paddle/fluid/tests/unittests/test_fleet_gradient_merge_meta_optimizer.py
...sts/unittests/test_fleet_gradient_merge_meta_optimizer.py
+13
-0
未找到文件。
python/paddle/distributed/fleet/meta_optimizers/amp_optimizer.py
浏览文件 @
ab049978
...
...
@@ -25,7 +25,6 @@ class AMPOptimizer(MetaOptimizerBase):
"LarsOptimizer"
,
"LambOptimizer"
,
"RecomputeOptimizer"
,
"GradientMergeOptimizer"
,
"GraphExecutionOptimizer"
,
]
self
.
meta_optimizers_black_list
=
[
"DGCOptimizer"
]
...
...
python/paddle/distributed/fleet/meta_optimizers/gradient_merge_optimizer.py
浏览文件 @
ab049978
...
...
@@ -21,6 +21,7 @@ class GradientMergeOptimizer(MetaOptimizerBase):
self
.
inner_opt
=
optimizer
self
.
wrapped_opt
=
None
self
.
meta_optimizers_white_list
=
[
"AMPOptimizer"
,
"LarsOptimizer"
,
"LambOptimizer"
,
"GraphExecutionOptimizer"
,
...
...
python/paddle/fluid/contrib/mixed_precision/decorator.py
浏览文件 @
ab049978
...
...
@@ -159,9 +159,6 @@ class OptimizerWithMixedPrecision(object):
params_grads
=
self
.
_optimizer
.
backward
(
self
.
_scaled_loss
,
startup_program
,
parameter_list
,
no_grad_set
,
callbacks
)
# Change the op_role_var attr for some ops, so that gradients
# transferred across GPUs can be FP16.
update_role_var_grad
(
train_program
,
params_grads
)
return
params_grads
def
apply_gradients
(
self
,
params_grads
):
...
...
@@ -176,6 +173,10 @@ class OptimizerWithMixedPrecision(object):
A list of optimize operators.
"""
# Change the op_role_var attr for some ops, so that gradients
# transferred across GPUs can be FP16.
update_role_var_grad
(
self
.
_train_program
,
params_grads
)
grads
=
[
g
for
_
,
g
in
params_grads
]
if
not
self
.
_is_distributed
:
with
self
.
_train_program
.
_optimized_guard
(
grads
):
...
...
python/paddle/fluid/tests/unittests/test_fleet_gradient_merge_meta_optimizer.py
浏览文件 @
ab049978
...
...
@@ -46,6 +46,19 @@ class TestFleetGradientMergeMetaOptimizer(TestFleetMetaOptimizer):
self
.
assertIn
(
'@GradientMerge'
,
''
.
join
(
vars
))
self
.
assertIn
(
'subprog'
,
''
.
join
(
vars
))
def
test_gm_amp_optimizer
(
self
):
train_prog
,
startup_prog
=
paddle
.
fluid
.
Program
(),
paddle
.
fluid
.
Program
(
)
avg_cost
,
strategy
=
self
.
net
(
train_prog
,
startup_prog
)
self
.
set_strategy
(
strategy
,
'gradient_merge'
)
self
.
set_strategy
(
strategy
,
'amp'
)
self
.
optimizer
(
avg_cost
,
strategy
,
train_prog
,
startup_prog
)
print
(
train_prog
)
vars
=
[
x
.
name
for
x
in
train_prog
.
list_vars
()]
self
.
assertIn
(
'@GradientMerge'
,
''
.
join
(
vars
))
self
.
assertIn
(
'cast'
,
''
.
join
(
vars
))
if
__name__
==
"__main__"
:
unittest
.
main
()
编辑
预览
Markdown
is supported
0%
请重试
或
添加新附件
.
添加附件
取消
You are about to add
0
people
to the discussion. Proceed with caution.
先完成此消息的编辑!
取消
想要评论请
注册
或
登录