Fix CUDA memory issues
This commit is contained in:
@@ -30,11 +30,14 @@ class SuperLinear(SuperModule):
|
||||
self._out_features = out_features
|
||||
self._bias = bias
|
||||
# weights to be optimized
|
||||
self._super_weight = torch.nn.Parameter(
|
||||
torch.Tensor(self.out_features, self.in_features)
|
||||
self.register_parameter(
|
||||
"_super_weight",
|
||||
torch.nn.Parameter(torch.Tensor(self.out_features, self.in_features)),
|
||||
)
|
||||
if self.bias:
|
||||
self._super_bias = torch.nn.Parameter(torch.Tensor(self.out_features))
|
||||
self.register_parameter(
|
||||
"_super_bias", torch.nn.Parameter(torch.Tensor(self.out_features))
|
||||
)
|
||||
else:
|
||||
self.register_parameter("_super_bias", None)
|
||||
self.reset_parameters()
|
||||
|
@@ -25,8 +25,8 @@ class SuperLayerNorm1D(SuperModule):
|
||||
self._eps = eps
|
||||
self._elementwise_affine = elementwise_affine
|
||||
if self._elementwise_affine:
|
||||
self.weight = nn.Parameter(torch.Tensor(self.in_dim))
|
||||
self.bias = nn.Parameter(torch.Tensor(self.in_dim))
|
||||
self.register_parameter("weight", nn.Parameter(torch.Tensor(self.in_dim)))
|
||||
self.register_parameter("bias", nn.Parameter(torch.Tensor(self.in_dim)))
|
||||
else:
|
||||
self.register_parameter("weight", None)
|
||||
self.register_parameter("bias", None)
|
||||
|
Reference in New Issue
Block a user