forked from gorgonia/gorgonia
-
Notifications
You must be signed in to change notification settings - Fork 0
/
values_cuda.go
74 lines (68 loc) · 1.84 KB
/
values_cuda.go
1
2
3
4
5
6
7
8
9
10
11
12
13
14
15
16
17
18
19
20
21
22
23
24
25
26
27
28
29
30
31
32
33
34
35
36
37
38
39
40
41
42
43
44
45
46
47
48
49
50
51
52
53
54
55
56
57
58
59
60
61
62
63
64
65
66
67
68
69
70
71
72
73
74
// +build cuda
package gorgonia
import (
"unsafe"
"github.com/chewxy/cu"
"github.com/chewxy/gorgonia/tensor"
"github.com/pkg/errors"
)
func valToDevicePointer(val Value) (mem cu.DevicePtr, err error) {
switch v := val.(type) {
case *tensor.Dense:
size := int64(v.DataSize() * int(v.Dtype().Size()))
if mem, err = cu.MemAlloc(size); err != nil {
err = errors.Wrapf(err, "Cannot get mem device pointer")
return
}
switch v.Dtype() {
case tensor.Float64:
data := v.Data().([]float64)
if err = cu.MemcpyHtoD(mem, unsafe.Pointer(&data[0]), size); err != nil {
err = errors.Wrapf(err, "Memcpy failed")
return
}
return
case tensor.Float32:
data := v.Data().([]float32)
if err = cu.MemcpyHtoD(mem, unsafe.Pointer(&data[0]), size); err != nil {
err = errors.Wrapf(err, "Memcpy failed")
return
}
return
default:
if err = cu.MemFree(mem); err != nil {
err = errors.Wrapf(err, "Unable to free mem properly")
return
}
return 0, errors.Errorf(unsupportedDtype, v.Dtype())
}
default:
}
return 0, errors.Errorf("Cannot convert %T to device pointer", val)
}
func devPtrToValue(val Value, mem cu.DevicePtr) (err error) {
switch v := val.(type) {
case *tensor.Dense:
size := int64(v.DataSize() * int(v.Dtype().Size()))
switch v.Dtype() {
case tensor.Float64:
data := v.Data().([]float64)
if err = cu.MemcpyDtoH(unsafe.Pointer(&data[0]), mem, size); err != nil {
err = errors.Wrapf(err, "Memcpy failed")
return
}
return nil
case tensor.Float32:
data := v.Data().([]float32)
if err = cu.MemcpyDtoH(unsafe.Pointer(&data[0]), mem, size); err != nil {
err = errors.Wrapf(err, "Memcpy failed")
return
}
return nil
default:
return errors.Errorf(unsupportedDtype, v.Dtype())
}
default:
}
return errors.Errorf("Cannot copy memory from device to %T", val)
}